{"data":{"node":{"slug":"floridsleeves-llmdebugger","name":"LLMDebugger","tagline":"A Large Language Model Debugger verifying runtime execution step by step","github_url":"https://github.com/FloridSleeves/LLMDebugger","owner":"FloridSleeves","repo":"LLMDebugger","owner_avatar_url":"https://avatars.githubusercontent.com/u/23695653?v=4","primary_language":"Python","stars":587,"forks":56,"topics":[],"archived":false,"github_pushed_at":"2024-09-10T23:32:12+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/floridsleeves-llmdebugger","markdown_url":"https://www.graphcanon.com/tools/floridsleeves-llmdebugger.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/floridsleeves-llmdebugger","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=floridsleeves-llmdebugger"},"categories":[{"slug":"developer-tools","name":"Developer Tools","url":"https://www.graphcanon.com/categories/developer-tools","markdown_url":"https://www.graphcanon.com/categories/developer-tools.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/developer-tools"},{"slug":"evaluation-observability","name":"Evaluation & Observability","url":"https://www.graphcanon.com/categories/evaluation-observability","markdown_url":"https://www.graphcanon.com/categories/evaluation-observability.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/evaluation-observability"}],"tags":[{"slug":"acl-24","name":"acl'24"},{"slug":"llm-debugging","name":"llm debugging"},{"slug":"python-debugger-for-ai","name":"python debugger for ai"},{"slug":"runtime-verification","name":"runtime verification"}],"edges":[],"neighbours":[{"slug":"confident-ai-deepeval","name":"deepeval","tagline":"LLM Evaluation Framework.","github_url":"https://github.com/confident-ai/deepeval","owner":"confident-ai","repo":"deepeval","owner_avatar_url":"https://avatars.githubusercontent.com/u/130858411?v=4","primary_language":"Python","stars":17226,"forks":1736,"topics":["evaluation-framework","evaluation-metrics","llm-evaluation","llm-evaluation-framework","llm-evaluation-metrics","python"],"archived":false,"github_pushed_at":"2026-07-27T11:33:31+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/confident-ai-deepeval","markdown_url":"https://www.graphcanon.com/tools/confident-ai-deepeval.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/confident-ai-deepeval","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=confident-ai-deepeval","shared_categories":["evaluation-observability"]},{"slug":"shishirpatil-gorilla","name":"gorilla","tagline":"Training and Evaluating LLMs for Function Calls (Tool Calls)","github_url":"https://github.com/ShishirPatil/gorilla","owner":"ShishirPatil","repo":"gorilla","owner_avatar_url":"https://avatars.githubusercontent.com/u/30296397?v=4","primary_language":"Python","stars":12988,"forks":1397,"topics":["api","api-documentation","chatgpt","claude-api","gpt-4-api","llm","openai-api","openai-functions"],"archived":false,"github_pushed_at":"2026-04-13T03:19:45+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/shishirpatil-gorilla","markdown_url":"https://www.graphcanon.com/tools/shishirpatil-gorilla.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/shishirpatil-gorilla","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=shishirpatil-gorilla","shared_categories":["evaluation-observability"]},{"slug":"zai-org-codegeex","name":"CodeGeeX","tagline":"CodeGeeX is an open multilingual code generation model implemented in Mindspore and available via PyTorch.","github_url":"https://github.com/zai-org/CodeGeeX","owner":"zai-org","repo":"CodeGeeX","owner_avatar_url":"https://avatars.githubusercontent.com/u/223098841?v=4","primary_language":"Python","stars":8809,"forks":688,"topics":["code-generation","pretrained-models","tools"],"archived":false,"github_pushed_at":"2024-08-13T05:59:38+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/zai-org-codegeex","markdown_url":"https://www.graphcanon.com/tools/zai-org-codegeex.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/zai-org-codegeex","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=zai-org-codegeex","shared_categories":[]},{"slug":"swe-bench-swe-bench","name":"SWE-bench","tagline":"Benchmark for assessing language models' capability to resolve real-world Github issues","github_url":"https://github.com/SWE-bench/SWE-bench","owner":"SWE-bench","repo":"SWE-bench","owner_avatar_url":"https://avatars.githubusercontent.com/u/139597579?v=4","primary_language":"Python","stars":5576,"forks":930,"topics":["benchmark","language-model","software-engineering"],"archived":false,"github_pushed_at":"2026-07-27T05:34:27+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/swe-bench-swe-bench","markdown_url":"https://www.graphcanon.com/tools/swe-bench-swe-bench.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/swe-bench-swe-bench","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=swe-bench-swe-bench","shared_categories":["evaluation-observability"]},{"slug":"eth-sri-lmql","name":"lmql","tagline":"A language for constraint-guided and efficient LLM programming.","github_url":"https://github.com/eth-sri/lmql","owner":"eth-sri","repo":"lmql","owner_avatar_url":"https://avatars.githubusercontent.com/u/5363413?v=4","primary_language":"Python","stars":4203,"forks":221,"topics":["chatgpt","huggingface","language-model","programming-language"],"archived":false,"github_pushed_at":"2025-05-22T07:32:31+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/eth-sri-lmql","markdown_url":"https://www.graphcanon.com/tools/eth-sri-lmql.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/eth-sri-lmql","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=eth-sri-lmql","shared_categories":[]},{"slug":"code-yeongyu-lazycodex","name":"lazycodex","tagline":"Agent harness for complex codebases with project memory and execution planning","github_url":"https://github.com/code-yeongyu/lazycodex","owner":"code-yeongyu","repo":"lazycodex","owner_avatar_url":"https://avatars.githubusercontent.com/u/11153873?v=4","primary_language":"TypeScript","stars":3189,"forks":198,"topics":["ai","ai-agents","claude","claude-code","cli","codex","developer-tools","lazy","lazycodex","oh-my-openagent","omo","openai","orchestration","typescript"],"archived":false,"github_pushed_at":"2026-08-09T08:51:31+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/code-yeongyu-lazycodex","markdown_url":"https://www.graphcanon.com/tools/code-yeongyu-lazycodex.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/code-yeongyu-lazycodex","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=code-yeongyu-lazycodex","shared_categories":["developer-tools"]},{"slug":"evalplus-evalplus","name":"evalplus","tagline":"Rigorous evaluation of LLM-synthesized code","github_url":"https://github.com/evalplus/evalplus","owner":"evalplus","repo":"evalplus","owner_avatar_url":"https://avatars.githubusercontent.com/u/132106461?v=4","primary_language":"Python","stars":1794,"forks":205,"topics":["benchmark","chatgpt","efficiency","gpt-4","large-language-models","program-synthesis","testing"],"archived":false,"github_pushed_at":"2025-10-02T22:56:38+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/evalplus-evalplus","markdown_url":"https://www.graphcanon.com/tools/evalplus-evalplus.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/evalplus-evalplus","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=evalplus-evalplus","shared_categories":["evaluation-observability"]},{"slug":"benman1-generative-ai-with-langchain","name":"generative_ai_with_langchain","tagline":"Build production-ready LLM applications and advanced agents using Python, LangChain, and LangGraph","github_url":"https://github.com/benman1/generative_ai_with_langchain","owner":"benman1","repo":"generative_ai_with_langchain","owner_avatar_url":"https://avatars.githubusercontent.com/u/10786684?v=4","primary_language":"Jupyter Notebook","stars":1400,"forks":582,"topics":["agent","chatgpt","claude","claude-3-5-sonnet","deepseek","deepseek-r1","gpt","gpt-4o","huggingface","langchain","langgraph","llamacpp","llms","ollama","openai"],"archived":false,"github_pushed_at":"2026-08-05T12:50:30+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/benman1-generative-ai-with-langchain","markdown_url":"https://www.graphcanon.com/tools/benman1-generative-ai-with-langchain.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/benman1-generative-ai-with-langchain","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=benman1-generative-ai-with-langchain","shared_categories":[]},{"slug":"bigcode-project-bigcode-evaluation-harness","name":"bigcode-evaluation-harness","tagline":"A framework for evaluating autoregressive code generation language models.","github_url":"https://github.com/bigcode-project/bigcode-evaluation-harness","owner":"bigcode-project","repo":"bigcode-evaluation-harness","owner_avatar_url":"https://avatars.githubusercontent.com/u/110470554?v=4","primary_language":"Python","stars":1055,"forks":261,"topics":[],"archived":false,"github_pushed_at":"2025-07-22T13:18:09+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/bigcode-project-bigcode-evaluation-harness","markdown_url":"https://www.graphcanon.com/tools/bigcode-project-bigcode-evaluation-harness.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/bigcode-project-bigcode-evaluation-harness","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=bigcode-project-bigcode-evaluation-harness","shared_categories":["evaluation-observability"]},{"slug":"livecodebench-livecodebench","name":"LiveCodeBench","tagline":"Holistic and contamination-free evaluation of large language models for code","github_url":"https://github.com/LiveCodeBench/LiveCodeBench","owner":"LiveCodeBench","repo":"LiveCodeBench","owner_avatar_url":"https://avatars.githubusercontent.com/u/161278213?v=4","primary_language":"Python","stars":925,"forks":195,"topics":["code-execution","code-generation","code-llms","code-repair","gpt-4","test-generation"],"archived":false,"github_pushed_at":"2025-07-16T00:58:38+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/livecodebench-livecodebench","markdown_url":"https://www.graphcanon.com/tools/livecodebench-livecodebench.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/livecodebench-livecodebench","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=livecodebench-livecodebench","shared_categories":["evaluation-observability"]},{"slug":"harleyszhang-llm-note","name":"llm_note","tagline":"LLM notes covering model inference transformer structures and framework analysis","github_url":"https://github.com/harleyszhang/llm_note","owner":"harleyszhang","repo":"llm_note","owner_avatar_url":"https://avatars.githubusercontent.com/u/37138671?v=4","primary_language":"Python","stars":888,"forks":90,"topics":["cuda-programming","kv-cache","llm","llm-inference","transformer-models","triton-kernels","vllm"],"archived":false,"github_pushed_at":"2026-08-19T06:46:41+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/harleyszhang-llm-note","markdown_url":"https://www.graphcanon.com/tools/harleyszhang-llm-note.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/harleyszhang-llm-note","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=harleyszhang-llm-note","shared_categories":[]},{"slug":"georgian-io-llm-finetuning-toolkit","name":"LLM-Finetuning-Toolkit","tagline":"Toolkit for fine-tuning and testing open-source large language models","github_url":"https://github.com/georgian-io/LLM-Finetuning-Toolkit","owner":"georgian-io","repo":"LLM-Finetuning-Toolkit","owner_avatar_url":"https://avatars.githubusercontent.com/u/10764713?v=4","primary_language":"Python","stars":870,"forks":107,"topics":["ablation-study","classification","falcon","fine-tuning","finetuning","flan-t5","large-language-models","llama2","llm-test","lora","mistral-7b","nlp","nlp-machine-learning","qlora","redpajama","summarization","unit-testing","zephyr"],"archived":false,"github_pushed_at":"2026-05-04T16:33:40+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/georgian-io-llm-finetuning-toolkit","markdown_url":"https://www.graphcanon.com/tools/georgian-io-llm-finetuning-toolkit.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/georgian-io-llm-finetuning-toolkit","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=georgian-io-llm-finetuning-toolkit","shared_categories":[]}]}}