{"library":"pdftext","type":"library","category":null,"description":"pdftext is a Python library designed for fast and accurate extraction of structured text from PDF documents. It focuses on efficiently parsing text, detecting elements like tables and links, and handling complex layouts. The current version is 0.6.3, and it's actively maintained with frequent minor releases addressing bug fixes and introducing new features.","language":"python","status":"active","version":"0.6.3","tags":["pdf","text-extraction","document-processing","nlp"],"install":[{"cmd":"pip install pdftext","imports":["import pdftext"]}],"homepage":null,"github":"https://github.com/VikParuchuri/pdftext","docs":null,"changelog":null,"pypi":"https://pypi.org/project/pdftext/","npm":null,"openapi_spec":null,"status_page":null,"smithery":null,"compatibility":{"summary":{"python_range":"3.10–3.9","success_rate":90,"avg_install_s":5.6,"avg_import_s":null,"wheel_type":"wheel"},"url":"https://checklist.day/v1/registry/pdftext/compatibility"},"provenance":{"verified_status":"passing","verified_at":"Tue Jun 30","last_verified":"Tue Jun 30","next_check":"Thu Jul 30","install_tag":null}}