{"library":"news-please","type":"library","category":null,"description":"news-please is an open-source, easy-to-use Python library designed for crawling news websites and extracting structured information from articles. It can recursively follow internal hyperlinks and read RSS feeds to fetch both recent and archived articles. The library also provides an API for programmatic use within Python applications and supports extracting articles from the commoncrawl.org news archive. It is currently active, with version 1.6.16 released, and maintains a regular release cadence.","language":"python","status":"active","version":"1.6.16","tags":["news","crawler","scraper","information extraction","web scraping","article extraction"],"install":[{"cmd":"pip install news-please","imports":["from newsplease import NewsPlease"]}],"homepage":null,"github":"https://github.com/fhamborg/news-please","docs":null,"changelog":null,"pypi":"https://pypi.org/project/news-please/","npm":null,"openapi_spec":null,"status_page":null,"smithery":null,"compatibility":{"summary":{"python_range":"3.10–3.9","success_rate":60,"avg_install_s":20.3,"avg_import_s":3.69,"wheel_type":"sdist"},"url":"https://checklist.day/v1/registry/news-please/compatibility"},"provenance":{"verified_status":"passing","verified_at":"Sun Jun 28","last_verified":"Sun Jun 28","next_check":"Tue Jul 28","install_tag":null}}