We are extracting the text from this PDF using the PdfPageBase.ExtractText() command but the return is being characters in format.
Could you check the attached PDF? Thanks.
PdfDocument doc = new PdfDocument();
// Read a pdf file
string output = @"..\..\data\Boletos vencimento 11.2022.pdf";
doc.LoadFromFile(output);
//Get the fonts used in PDF
PdfUsedFont[] fonts = doc.UsedFonts;
//Travel the fonts used in PDF
foreach (PdfUsedFont font in fonts)
{
if(font.Name == "F1"||font.Name=="F2"||font.Name=="F3")
{
Console.WriteLine("Sorry,this pdf document does not support text extraction ");
break;
}
}