{"data":{"node":{"slug":"urchade-gliner","name":"GLiNER","tagline":"Generalist and Lightweight Model for Named Entity Recognition","github_url":"https://github.com/urchade/GLiNER","owner":"urchade","repo":"GLiNER","owner_avatar_url":"https://avatars.githubusercontent.com/u/38214774?v=4","primary_language":"Python","stars":3545,"forks":299,"topics":["information-extraction","large-language-models","named-entity-recognition","natural-language-processing","prompt-tuning"],"archived":false,"github_pushed_at":"2026-08-10T09:21:43+00:00","maintenance_label":"Active","stars_delta_30d":143,"url":"https://www.graphcanon.com/tools/urchade-gliner","markdown_url":"https://www.graphcanon.com/tools/urchade-gliner.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/urchade-gliner","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=urchade-gliner"},"categories":[{"slug":"data-retrieval","name":"Data & Retrieval","url":"https://www.graphcanon.com/categories/data-retrieval","markdown_url":"https://www.graphcanon.com/categories/data-retrieval.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/data-retrieval"},{"slug":"model-training","name":"Model Training","url":"https://www.graphcanon.com/categories/model-training","markdown_url":"https://www.graphcanon.com/categories/model-training.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/model-training"}],"tags":[{"slug":"information-extraction","name":"information-extraction"},{"slug":"large-language-models","name":"large language models"},{"slug":"named-entity-recognition","name":"named-entity-recognition"},{"slug":"natural-language-processing","name":"natural-language-processing"},{"slug":"prompt-tuning","name":"prompt-tuning"}],"edges":[{"type":"depends_on","direction":"out","explanation":"The repository focuses on NER, and huggingface's transformers is a widely used framework for implementing such models.","successor_context":null,"tool":{"slug":"huggingface-transformers","name":"transformers","tagline":"Transformers: the model-definition framework for state-of-the-art machine learning models in text, vision, audio, and multimodal models","github_url":"https://github.com/huggingface/transformers","owner":"huggingface","repo":"transformers","owner_avatar_url":"https://avatars.githubusercontent.com/u/25720743?v=4","primary_language":"Python","stars":164121,"forks":34249,"topics":["audio","deep-learning","deepseek","gemma","glm","hacktoberfest","llm","machine-learning","model-hub","natural-language-processing","nlp","pretrained-models","python","pytorch","pytorch-transformers","qwen","speech-recognition","transformer","vlm"],"archived":false,"github_pushed_at":"2026-08-15T22:28:12+00:00","maintenance_label":"Very active","stars_delta_30d":1457,"url":"https://www.graphcanon.com/tools/huggingface-transformers","markdown_url":"https://www.graphcanon.com/tools/huggingface-transformers.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/huggingface-transformers","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=huggingface-transformers"}},{"type":"integrates_with","direction":"out","explanation":"GLiNER relies on transformer models for Named Entity Recognition tasks; therefore, it integrates well with Hugging Face's transformers library which provides a wide range of pre-trained transformer models.","successor_context":null,"tool":{"slug":"huggingface-transformers","name":"transformers","tagline":"Transformers: the model-definition framework for state-of-the-art machine learning models in text, vision, audio, and multimodal models","github_url":"https://github.com/huggingface/transformers","owner":"huggingface","repo":"transformers","owner_avatar_url":"https://avatars.githubusercontent.com/u/25720743?v=4","primary_language":"Python","stars":164121,"forks":34249,"topics":["audio","deep-learning","deepseek","gemma","glm","hacktoberfest","llm","machine-learning","model-hub","natural-language-processing","nlp","pretrained-models","python","pytorch","pytorch-transformers","qwen","speech-recognition","transformer","vlm"],"archived":false,"github_pushed_at":"2026-08-15T22:28:12+00:00","maintenance_label":"Very active","stars_delta_30d":1457,"url":"https://www.graphcanon.com/tools/huggingface-transformers","markdown_url":"https://www.graphcanon.com/tools/huggingface-transformers.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/huggingface-transformers","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=huggingface-transformers"}},{"type":"integrates_with","direction":"out","explanation":"GLiNER performs named entity recognition which can be integrated with langextract to extract structured information from unstructured text using LLMs.","successor_context":null,"tool":{"slug":"google-langextract","name":"langextract","tagline":"A Python library for extracting structured information from unstructured text using LLMs.","github_url":"https://github.com/google/langextract","owner":"google","repo":"langextract","owner_avatar_url":"https://avatars.githubusercontent.com/u/1342004?v=4","primary_language":"Python","stars":38400,"forks":2693,"topics":["gemini","gemini-ai","gemini-api","gemini-flash","gemini-pro","information-extration","large-language-models","llm","nlp","python","structured-data"],"archived":false,"github_pushed_at":"2026-08-11T15:31:39+00:00","maintenance_label":"Very active","stars_delta_30d":1241,"url":"https://www.graphcanon.com/tools/google-langextract","markdown_url":"https://www.graphcanon.com/tools/google-langextract.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/google-langextract","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=google-langextract"}}],"neighbours":[{"slug":"lightning-ai-litgpt","name":"litgpt","tagline":"High-performance LLMs with recipes for pretraining, finetuning and deployment","github_url":"https://github.com/Lightning-AI/litgpt","owner":"Lightning-AI","repo":"litgpt","owner_avatar_url":"https://avatars.githubusercontent.com/u/58386951?v=4","primary_language":"Python","stars":13605,"forks":1483,"topics":["ai","artificial-intelligence","deep-learning","large-language-models","llm","llm-inference","llms"],"archived":false,"github_pushed_at":"2026-07-20T10:24:12+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/lightning-ai-litgpt","markdown_url":"https://www.graphcanon.com/tools/lightning-ai-litgpt.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/lightning-ai-litgpt","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=lightning-ai-litgpt","shared_categories":["model-training"]},{"slug":"eugeneyan-open-llms","name":"open-llms","tagline":"A list of open LLMs available for commercial use.","github_url":"https://github.com/eugeneyan/open-llms","owner":"eugeneyan","repo":"open-llms","owner_avatar_url":"https://avatars.githubusercontent.com/u/6831355?v=4","primary_language":null,"stars":12849,"forks":985,"topics":["commercial","large-language-models","llm","llms"],"archived":false,"github_pushed_at":"2025-02-13T06:37:12+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/eugeneyan-open-llms","markdown_url":"https://www.graphcanon.com/tools/eugeneyan-open-llms.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/eugeneyan-open-llms","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=eugeneyan-open-llms","shared_categories":[]},{"slug":"huggingface-tokenizers","name":"tokenizers","tagline":"💥 Fast State-of-the-Art Tokenizers optimized for Research and Production","github_url":"https://github.com/huggingface/tokenizers","owner":"huggingface","repo":"tokenizers","owner_avatar_url":"https://avatars.githubusercontent.com/u/25720743?v=4","primary_language":"Rust","stars":10940,"forks":1160,"topics":["bert","gpt","language-model","natural-language-processing","natural-language-understanding","nlp","transformers"],"archived":false,"github_pushed_at":"2026-08-01T12:35:36+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/huggingface-tokenizers","markdown_url":"https://www.graphcanon.com/tools/huggingface-tokenizers.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/huggingface-tokenizers","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=huggingface-tokenizers","shared_categories":["model-training"]},{"slug":"data-privacy-stack-presidio","name":"presidio","tagline":"A framework for detecting and anonymizing sensitive data","github_url":"https://github.com/data-privacy-stack/presidio","owner":"data-privacy-stack","repo":"presidio","owner_avatar_url":"https://avatars.githubusercontent.com/u/275623515?v=4","primary_language":"Python","stars":10395,"forks":1237,"topics":["anonymization","data-anonymization","data-masking","data-obfuscation","data-privacy","data-redaction","de-identification","guardrails","image-redactor","named-entity-recognition","nlp","personally-identifiable-information","phi","pii","pii-detection","privacy","python","sensitive-data","spacy","transformers"],"archived":false,"github_pushed_at":"2026-08-08T21:25:09+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/data-privacy-stack-presidio","markdown_url":"https://www.graphcanon.com/tools/data-privacy-stack-presidio.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/data-privacy-stack-presidio","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=data-privacy-stack-presidio","shared_categories":["data-retrieval"]},{"slug":"fareedkhan-dev-train-llm-from-scratch","name":"train-llm-from-scratch","tagline":"A straightforward method for training your LLM from raw text to aligned model generation","github_url":"https://github.com/FareedKhan-dev/train-llm-from-scratch","owner":"FareedKhan-dev","repo":"train-llm-from-scratch","owner_avatar_url":"https://avatars.githubusercontent.com/u/63067900?v=4","primary_language":"Python","stars":9141,"forks":1264,"topics":["gemini","large-language-models","llm","openai","training","transformers"],"archived":false,"github_pushed_at":"2026-08-17T05:07:26+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/fareedkhan-dev-train-llm-from-scratch","markdown_url":"https://www.graphcanon.com/tools/fareedkhan-dev-train-llm-from-scratch.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/fareedkhan-dev-train-llm-from-scratch","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=fareedkhan-dev-train-llm-from-scratch","shared_categories":["model-training"]},{"slug":"wangrongsheng-awesome-llm-resources","name":"awesome-LLM-resources","tagline":"Summary of the world's best LLM resources.","github_url":"https://github.com/WangRongsheng/awesome-LLM-resources","owner":"WangRongsheng","repo":"awesome-LLM-resources","owner_avatar_url":"https://avatars.githubusercontent.com/u/55651568?v=4","primary_language":null,"stars":8845,"forks":950,"topics":["awesome-list","book","course","large-language-models","llama","llm","mistral","openai","qwen","rag","retrieval-augmented-generation","webui"],"archived":false,"github_pushed_at":"2026-08-14T15:54:28+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/wangrongsheng-awesome-llm-resources","markdown_url":"https://www.graphcanon.com/tools/wangrongsheng-awesome-llm-resources.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/wangrongsheng-awesome-llm-resources","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=wangrongsheng-awesome-llm-resources","shared_categories":["model-training"]},{"slug":"datajuicer-data-juicer","name":"data-juicer","tagline":"Data processing for and with foundation models","github_url":"https://github.com/datajuicer/data-juicer","owner":"datajuicer","repo":"data-juicer","owner_avatar_url":"https://avatars.githubusercontent.com/u/223222708?v=4","primary_language":"Python","stars":6897,"forks":404,"topics":["data","data-analysis","data-pipeline","data-processing","data-science","data-visualization","foundation-models","instruction-tuning","large-language-models","llm","llms","multi-modal","pre-training","synthetic-data"],"archived":false,"github_pushed_at":"2026-08-13T09:19:31+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/datajuicer-data-juicer","markdown_url":"https://www.graphcanon.com/tools/datajuicer-data-juicer.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/datajuicer-data-juicer","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=datajuicer-data-juicer","shared_categories":["model-training","data-retrieval"]},{"slug":"xlite-dev-awesome-llm-inference","name":"Awesome-LLM-Inference","tagline":"A curated list of LLM/VLM inference papers with codes","github_url":"https://github.com/xlite-dev/Awesome-LLM-Inference","owner":"xlite-dev","repo":"Awesome-LLM-Inference","owner_avatar_url":"https://avatars.githubusercontent.com/u/204302598?v=4","primary_language":"Python","stars":5415,"forks":428,"topics":["awesome-llm","deepseek","deepseek-r1","deepseek-v3","flash-attention","flash-attention-3","flash-mla","llm-inference","minimax-01","mla","paged-attention","qwen3","tensorrt-llm","vllm"],"archived":false,"github_pushed_at":"2026-06-23T03:48:43+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/xlite-dev-awesome-llm-inference","markdown_url":"https://www.graphcanon.com/tools/xlite-dev-awesome-llm-inference.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/xlite-dev-awesome-llm-inference","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=xlite-dev-awesome-llm-inference","shared_categories":[]},{"slug":"maziyarpanahi-openmed","name":"openmed","tagline":"Local-first healthcare AI for clinical NER and HIPAA PII de-identification.","github_url":"https://github.com/maziyarpanahi/openmed","owner":"maziyarpanahi","repo":"openmed","owner_avatar_url":"https://avatars.githubusercontent.com/u/5762953?v=4","primary_language":"Python","stars":4974,"forks":619,"topics":["android","clinical-nlp","healthcare","hipaa","ios","javascript","llm","local-llm","mlx","ner","nlp","on-device","on-premise","pii","pii-detection","python","skills","sovereign-ai","swift","swiftui"],"archived":false,"github_pushed_at":"2026-08-12T14:28:04+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/maziyarpanahi-openmed","markdown_url":"https://www.graphcanon.com/tools/maziyarpanahi-openmed.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/maziyarpanahi-openmed","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=maziyarpanahi-openmed","shared_categories":[]},{"slug":"dbiir-uer-py","name":"UER-py","tagline":"Open Source Pre-training Model Framework in PyTorch & Pre-trained Model Zoo","github_url":"https://github.com/dbiir/UER-py","owner":"dbiir","repo":"UER-py","owner_avatar_url":"https://avatars.githubusercontent.com/u/13671736?v=4","primary_language":"Python","stars":3110,"forks":520,"topics":["albert","bart","bert","chinese","classification","clue","elmo","fine-tuning","gpt","gpt-2","model-zoo","natural-language-processing","ner","pegasus","pre-training","pytorch","roberta","t5","unilm","xlm-roberta"],"archived":false,"github_pushed_at":"2024-05-09T11:12:55+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/dbiir-uer-py","markdown_url":"https://www.graphcanon.com/tools/dbiir-uer-py.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/dbiir-uer-py","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=dbiir-uer-py","shared_categories":["model-training"]},{"slug":"huangowen-awesome-llm-compression","name":"Awesome-LLM-Compression","tagline":"Awesome LLM compression research papers and tools to accelerate LLM training and inference.","github_url":"https://github.com/HuangOwen/Awesome-LLM-Compression","owner":"HuangOwen","repo":"Awesome-LLM-Compression","owner_avatar_url":"https://avatars.githubusercontent.com/u/24937399?v=4","primary_language":null,"stars":1859,"forks":129,"topics":[],"archived":false,"github_pushed_at":"2026-06-30T15:26:46+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/huangowen-awesome-llm-compression","markdown_url":"https://www.graphcanon.com/tools/huangowen-awesome-llm-compression.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/huangowen-awesome-llm-compression","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=huangowen-awesome-llm-compression","shared_categories":[]},{"slug":"tencent-tencentpretrain","name":"TencentPretrain","tagline":"Tencent Pre-training framework in PyTorch & Pre-trained Model Zoo","github_url":"https://github.com/Tencent/TencentPretrain","owner":"Tencent","repo":"TencentPretrain","owner_avatar_url":"https://avatars.githubusercontent.com/u/18461506?v=4","primary_language":"Python","stars":1090,"forks":148,"topics":["albert","bart","bert","chinese","classification","clue","elmo","fine-tuning","gpt","gpt-2","model-zoo","natural-language-processing","ner","pegasus","pre-training","pytorch","roberta","t5","unilm","xlm-roberta"],"archived":false,"github_pushed_at":"2024-08-04T11:53:43+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/tencent-tencentpretrain","markdown_url":"https://www.graphcanon.com/tools/tencent-tencentpretrain.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/tencent-tencentpretrain","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=tencent-tencentpretrain","shared_categories":["model-training"]}]}}