+
    Žvj`  ã                   óþ   € R t ^ RIt^ RIt^ RIHtHt ^ RIHt ^ RIH	t	H
t
Ht ^ RIHtHtHtHtHt ^ RIHt ^ RIHtHt ^ RIHtHt ^ R	IHt ^ R
IHtHt ^ RIH t  ^ RI!H"t"H#t#H$t$ RR R llt%RR R llt&RR R llt'R# )zIFunctions that can be used for the most common use-cases for pdfminer.sixN)Ú	ContainerÚIterator)ÚStringIO)ÚAnyÚBinaryIOÚcast)ÚHOCRConverterÚHTMLConverterÚPDFPageAggregatorÚTextConverterÚXMLConverter)ÚImageWriter)ÚLAParamsÚLTPage)Ú	PDFDeviceÚTagExtractor)ÚPDFValueError)ÚPDFPageInterpreterÚPDFResourceManager)ÚPDFPage)ÚAnyIOÚ
FileOrNameÚopen_filenamec          "      ó  € V ^8„  d   QhR\         R\        R\        R\        R\        R,          R\        R\
        \        ,          R,          R	\        R
\        R\        R\        R\        R,          R\        R\        R\        R\        RR/# )é   ÚinfÚoutfpÚoutput_typeÚcodecÚlaparamsNÚmaxpagesÚpage_numbersÚpasswordÚscaleÚrotationÚ
layoutmodeÚ
output_dirÚstrip_controlÚdebugÚdisable_cachingÚkwargsÚreturn)	r   r   Ústrr   Úintr   ÚfloatÚboolr   )Úformats   "ÚW/var/www/html/daimler-pipeline/venv/lib/python3.14/site-packages/pdfminer/high_level.pyÚ__annotate__r2      sÖ   € ÷ wñ wÜ	ðwäðwô ðwô ð	wô
 ˜�oðwô ðwô œC•. 4Õ'ðwô ðwô ðwô ðwô ðwô �d•
ðwô ðwô ðwô ðwô  ð!wð" 
ñ#wó    c           
     óV  € V'       d3   \         P                  ! 4       P                  \         P                  4       RpV'       d   \	        V4      p\        V'       * R7      pRpVR8w  d0   V\        P                  8X  d   \        P                  P                  pVR8X  d   \        VVVVVR7      pMVR8X  d   \        VVVVVVR7      pMfVR8X  d   \        VVVVV
VVR7      pMLVR	8X  d   \        VVVVVR
7      pM4VR8X  d   \        V\        \        V4      VR7      pMRV 2p\!        V4      hVf   Q h\#        VV4      p\$        P&                  ! V VVVV'       * R7       F3  pVP(                  V	,           R,          Vn        VP+                  V4       K5  	  VP-                  4        R# )a  Parses text from inf-file and writes to outfp file-like object.

Takes loads of optional arguments but the defaults are somewhat sane.
Beware laparams: Including an empty LAParams is not the same as passing
None!

:param inf: a file-like object to read PDF structure from, such as a
    file handler (using the builtin `open()` function) or a `BytesIO`.
:param outfp: a file-like object to write the text to.
:param output_type: May be 'text', 'xml', 'html', 'hocr', 'tag'.
    Only 'text' works properly.
:param codec: Text decoding codec
:param laparams: An LAParams object from pdfminer.layout. Default is None
    but may not layout correctly.
:param maxpages: How many pages to stop parsing after
:param page_numbers: zero-indexed page numbers to operate on.
:param password: For encrypted PDFs, the password to decrypt.
:param scale: Scale factor
:param rotation: Rotation factor
:param layoutmode: Default is 'normal', see
    pdfminer.converter.HTMLConverter
:param output_dir: If given, creates an ImageWriter for extracted images.
:param strip_control: Does what it says on the tin
:param debug: Output more logging data
:param disable_caching: Does what it says on the tin
:param other:
:return: nothing, acting as it does on two streams. Use StringIO to get
    strings.
N©ÚcachingÚtext)r   r   ÚimagewriterÚxml)r   r   r8   ÚstripcontrolÚhtml)r   r#   r%   r   r8   Úhocr)r   r   r:   Útag)r   z1Output type can be text, html, xml or tag but is ©r    r"   r6   ih  )ÚloggingÚ	getLoggerÚsetLevelÚDEBUGr   r   ÚsysÚstdoutÚbufferr   r   r	   r   r   r   r   r   r   r   Ú	get_pagesÚrotateÚprocess_pageÚclose)r   r   r   r   r   r    r!   r"   r#   r$   r%   r&   r'   r(   r)   r*   r8   ÚrsrcmgrÚdeviceÚmsgÚinterpreterÚpages   &&&&&&&&&&&&&&&,      r1   Úextract_text_to_fprO      s¤  € ÷^ Ü×ÒÓ×$Ñ$¤W§]¡]Ô3à€KßÜ! *Ó-ˆä ¨_Ô)<Ô=€GØ#€Fà�fÔ ¬#¯*©*Ô!4Ü—
‘
×!Ñ!ˆà�fÔÜØØØØØ#ô
‰ð 
˜Ô	ÜØØØØØ#Ø&ô
‰ð 
˜Ô	ÜØØØØØ!ØØ#ô
‰ð 
˜Ô	ÜØØØØØ&ô
‰ð 
˜Ô	ä˜g¤t¬H°eÓ'<ÀEÔJ‰ð BÀ+ÀÐOˆÜ˜CÓ Ð àÒÐÐÜ$ W¨fÓ5€KÜ×!Ò!ØØØØØ#Ô#÷ˆð —{‘{ XÕ-°Õ4ˆŒØ× Ñ  Ö&ñð ‡L�L†Nr3   c                óª   € V ^8„  d   QhR\         R\        R\        \        ,          R,          R\        R\        R\        R\
        R,          R	\        /# )
r   Úpdf_filer"   r!   Nr    r6   r   r   r+   )r   r,   r   r-   r/   r   )r0   s   "r1   r2   r2   “   se   € ÷ ((ñ ((Üð((äð((ô œC•. 4Õ'ð((ô ð	((ô
 ð((ô ð((ô ˜�oð((ô 	ñ((r3   c                óæ  € Vf   \        4       p\        V R4      ;_uu_ 4       p\        4       ;_uu_ 4       p\        \        V4      p\        VR7      p	\        W˜WVR7      p
\        Wš4      p\        P                  ! VVVVVR7       F  pVP                  V4       K  	  VP                  4       uuRRR4       uuRRR4       #   + '       g   i     M; iRRR4       R#   + '       g   i     R# ; i)aK  Parse and return the text contained in a PDF file.

:param pdf_file: Either a file path or a file-like object for the PDF file
    to be worked on.
:param password: For encrypted PDFs, the password to decrypt.
:param page_numbers: List of zero-indexed page numbers to extract.
:param maxpages: The maximum number of pages to parse
:param caching: If resources should be cached
:param codec: Text decoding codec
:param laparams: An LAParams object from pdfminer.layout. If None, uses
    some default settings that often work well.
:return: a string containing all of the text extracted.
NÚrbr5   )r   r   r>   )r   r   r   r   r   r   r   r   r   rF   rH   Úgetvalue)rQ   r"   r!   r    r6   r   r   ÚfpÚoutput_stringrJ   rK   rM   rN   s   &&&&&&&      r1   Úextract_textrW   “   sµ   € ð, ÒÜ“:ˆä	�x ×	&Ô	&¨"¬h¯j¬j¸MÜ”(˜BÓˆÜ$¨WÔ5ˆÜ˜w¸UÔVˆÜ(¨Ó9ˆä×%Ò%ØØØØØ÷
ˆDð ×$Ñ$ TÖ*ñ
ð ×%Ñ%Ó'÷ /9©j×	&Ò	&¯j¯jú×	&×	&×	&Ò	&ús#   £C¶A7C	Â-
CÃCÃCÃC0	c                ó´   € V ^8„  d   QhR\         R\        R\        \        ,          R,          R\        R\        R\
        R,          R\        \        ,          /# )	r   rQ   r"   r!   Nr    r6   r   r+   )r   r,   r   r-   r/   r   r   r   )r0   s   "r1   r2   r2   ¾   s`   € ÷ %ñ %Üð%äð%ô œC•. 4Õ'ð%ô ð	%ô
 ð%ô ˜�oð%ô ŒfÕñ%r3   c           
   #  ó€  "  € Vf   \        4       p\        V R4      ;_uu_ 4       p\        \        V4      p\	        VR7      p\        WuR7      p\        Wx4      p	\        P                  ! VVVVVR7       F(  p
V	P                  V
4       VP                  4       pVx € K*  	  RRR4       R#   + '       g   i     R# ; i5i)a÷  Extract and yield LTPage objects

:param pdf_file: Either a file path or a file-like object for the PDF file
    to be worked on.
:param password: For encrypted PDFs, the password to decrypt.
:param page_numbers: List of zero-indexed page numbers to extract.
:param maxpages: The maximum number of pages to parse
:param caching: If resources should be cached
:param laparams: An LAParams object from pdfminer.layout. If None, uses
    some default settings that often work well.
:return: LTPage objects
NrS   r5   )r   r>   )r   r   r   r   r   r
   r   r   rF   rH   Ú
get_result)rQ   r"   r!   r    r6   r   rU   Úresource_managerrK   rM   rN   Úlayouts   &&&&&&      r1   Úextract_pagesr]   ¾   s¥   é € ð( ÒÜ“:ˆä	�x ×	&Ô	&¨"Ü”(˜BÓˆÜ-°gÔ>ÐÜ"Ð#3ÔGˆÜ(Ð)9ÓBˆÜ×%Ò%ØØØØØ÷
ˆDð ×$Ñ$ TÔ*Ø×&Ñ&Ó(ˆFØŒLñ
÷ 
'×	&×	&Ò	&üs   ‚#B>¥A;B*Â 
B>Â*B;	Â5	B>)r7   úutf-8Né    NÚ g      ð?r_   ÚnormalNFFF)r`   Nr_   Tr^   N)r`   Nr_   TN)(Ú__doc__r?   rC   Úcollections.abcr   r   Úior   Útypingr   r   r   Úpdfminer.converterr   r	   r
   r   r   Úpdfminer.imager   Úpdfminer.layoutr   r   Úpdfminer.pdfdevicer   r   Úpdfminer.pdfexceptionsr   Úpdfminer.pdfinterpr   r   Úpdfminer.pdfpager   Úpdfminer.utilsr   r   r   rO   rW   r]   © r3   r1   Ú<module>ro      sU   ðÙ Oã Û 
ß /Ý ß &Ñ &÷õ õ 'ß ,ß 6Ý 0ß EÝ $ß ;Ñ ;÷w÷t((÷V%ñ %r3   