{"library":"segments","type":"library","category":null,"description":"Segments provides functions to tokenize and segment strings of text into individual characters or graphemes, and into segments according to orthography profiles. It is particularly useful for linguistic data processing using CLDF (Cross-Linguistic Data Formats). The library typically sees a few releases per year, with major versions introducing updates to Unicode standards.","language":"python","status":"active","version":"2.4.0","tags":["text processing","linguistics","unicode","segmentation","tokenization","graphemes","cldf"],"install":[{"cmd":"pip install segments","imports":["from segments import tokenize","from segments import Tokenizer","from segments import Profile"]}],"homepage":null,"github":"https://github.com/cldf/segments","docs":null,"changelog":null,"pypi":"https://pypi.org/project/segments/","npm":null,"openapi_spec":null,"status_page":null,"smithery":null,"compatibility":{"summary":{"python_range":"3.10–3.9","success_rate":100,"avg_install_s":4.4,"avg_import_s":null,"wheel_type":"wheel"},"url":"https://checklist.day/v1/registry/segments/compatibility"},"provenance":{"verified_status":"import_fail","verified_at":"Fri Jul 03","last_verified":"Fri Jul 03","next_check":"Fri Jul 10","install_tag":null}}