import re

with open(r"assets\pdf\jnanadeepa-may-aug-2025.pdf", "rb") as f:
    raw = f.read()

# find readable strings
strings = re.findall(b'[\x20-\x7E]{4,}', raw)
for s in strings[:40]:
    try:
        decoded = s.decode('utf-8', errors='ignore')
        if any(w in decoded.lower() for w in ['jnanadeepa', 'pope', 'francis', 'author', 'father', 'tor', 'vol']):
            print("Found:", decoded)
    except:
        pass
