{"data":{"node":{"slug":"evolvinglmms-lab-lmms-eval","name":"lmms-eval","tagline":"One-for-All Multimodal Evaluation Toolkit Across Text, Image, Video, and Audio Tasks","github_url":"https://github.com/EvolvingLMMs-Lab/lmms-eval","owner":"EvolvingLMMs-Lab","repo":"lmms-eval","owner_avatar_url":"https://avatars.githubusercontent.com/u/154951679?v=4","primary_language":"Python","stars":4368,"forks":639,"topics":["agi","audio-evaluation","benchmark","evaluation","large-language-models","llm-evaluation","multimodal","multimodal-evaluation","video-understanding","vision-language-model","vlm"],"archived":false,"github_pushed_at":"2026-08-06T02:22:23+00:00","maintenance_label":"Active","stars_delta_30d":52,"url":"https://www.graphcanon.com/tools/evolvinglmms-lab-lmms-eval","markdown_url":"https://www.graphcanon.com/tools/evolvinglmms-lab-lmms-eval.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/evolvinglmms-lab-lmms-eval","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=evolvinglmms-lab-lmms-eval"},"categories":[{"slug":"evaluation-observability","name":"Evaluation & Observability","url":"https://www.graphcanon.com/categories/evaluation-observability","markdown_url":"https://www.graphcanon.com/categories/evaluation-observability.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/evaluation-observability"}],"tags":[{"slug":"agi","name":"agi"},{"slug":"audio-evaluation","name":"audio-evaluation"},{"slug":"benchmark","name":"benchmark"},{"slug":"evaluation","name":"evaluation"},{"slug":"large-language-models","name":"large language models"},{"slug":"llm-evaluation","name":"llm-evaluation"},{"slug":"multimodal","name":"multimodal"},{"slug":"multimodal-evaluation","name":"multimodal-evaluation"}],"edges":[{"type":"integrates_with","direction":"out","explanation":"Both repositories provide tools for evaluating and testing generative AI systems, with lmms-eval focusing on multimodal evaluations while Giskard OSS focuses broadly on evals and red teaming.","successor_context":null,"tool":{"slug":"giskard-ai-giskard-oss","name":"giskard-oss","tagline":"Open-Source Evaluation & Testing library for LLM Agents","github_url":"https://github.com/Giskard-AI/giskard-oss","owner":"Giskard-AI","repo":"giskard-oss","owner_avatar_url":"https://avatars.githubusercontent.com/u/71782571?v=4","primary_language":"Python","stars":5727,"forks":511,"topics":["agent-evaluation","ai-red-team","ai-security","ai-testing","fairness-ai","llm","llm-eval","llm-evaluation","llm-security","llmops","ml-testing","ml-validation","mlops","rag-evaluation","red-team-tools","responsible-ai","trustworthy-ai"],"archived":false,"github_pushed_at":"2026-08-01T23:22:37+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/giskard-ai-giskard-oss","markdown_url":"https://www.graphcanon.com/tools/giskard-ai-giskard-oss.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/giskard-ai-giskard-oss","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=giskard-ai-giskard-oss"}},{"type":"alternative","direction":"out","explanation":"Both repositories deal with multimodal large language models, but they approach the evaluation and listing of these models differently.","successor_context":null,"tool":{"slug":"bradyfu-awesome-multimodal-large-language-models","name":"Awesome-Multimodal-Large-Language-Models","tagline":"Latest Advances on Multimodal Large Language Models","github_url":"https://github.com/BradyFU/Awesome-Multimodal-Large-Language-Models","owner":"BradyFU","repo":"Awesome-Multimodal-Large-Language-Models","owner_avatar_url":"https://avatars.githubusercontent.com/u/54254631?v=4","primary_language":null,"stars":17978,"forks":1133,"topics":["chain-of-thought","in-context-learning","instruction-following","instruction-tuning","large-language-models","large-vision-language-model","large-vision-language-models","multi-modality","multimodal-chain-of-thought","multimodal-in-context-learning","multimodal-instruction-tuning","multimodal-large-language-models","visual-instruction-tuning"],"archived":false,"github_pushed_at":"2026-08-14T17:17:50+00:00","maintenance_label":"Very active","stars_delta_30d":29,"url":"https://www.graphcanon.com/tools/bradyfu-awesome-multimodal-large-language-models","markdown_url":"https://www.graphcanon.com/tools/bradyfu-awesome-multimodal-large-language-models.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/bradyfu-awesome-multimodal-large-language-models","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=bradyfu-awesome-multimodal-large-language-models"}},{"type":"related","direction":"out","explanation":"The repository provides a list of resources regarding multmodal LLMs, which is the domain that lmms-eval aims to evaluate and improve upon, though it does not directly integrate or depend on this resource.","successor_context":null,"tool":{"slug":"bradyfu-awesome-multimodal-large-language-models","name":"Awesome-Multimodal-Large-Language-Models","tagline":"Latest Advances on Multimodal Large Language Models","github_url":"https://github.com/BradyFU/Awesome-Multimodal-Large-Language-Models","owner":"BradyFU","repo":"Awesome-Multimodal-Large-Language-Models","owner_avatar_url":"https://avatars.githubusercontent.com/u/54254631?v=4","primary_language":null,"stars":17978,"forks":1133,"topics":["chain-of-thought","in-context-learning","instruction-following","instruction-tuning","large-language-models","large-vision-language-model","large-vision-language-models","multi-modality","multimodal-chain-of-thought","multimodal-in-context-learning","multimodal-instruction-tuning","multimodal-large-language-models","visual-instruction-tuning"],"archived":false,"github_pushed_at":"2026-08-14T17:17:50+00:00","maintenance_label":"Very active","stars_delta_30d":29,"url":"https://www.graphcanon.com/tools/bradyfu-awesome-multimodal-large-language-models","markdown_url":"https://www.graphcanon.com/tools/bradyfu-awesome-multimodal-large-language-models.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/bradyfu-awesome-multimodal-large-language-models","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=bradyfu-awesome-multimodal-large-language-models"}},{"type":"alternative","direction":"out","explanation":"Both tools are designed to evaluate LLMs by providing ways to test and red-team LLM applications, but they use different methodologies for evaluation.","successor_context":null,"tool":{"slug":"promptfoo-promptfoo","name":"promptfoo","tagline":"Tool for evaluating prompts and AI agents by comparing performance across various models and red teaming.","github_url":"https://github.com/promptfoo/promptfoo","owner":"promptfoo","repo":"promptfoo","owner_avatar_url":"https://avatars.githubusercontent.com/u/137907881?v=4","primary_language":"TypeScript","stars":23838,"forks":2147,"topics":["ci","ci-cd","cicd","evaluation","evaluation-framework","llm","llm-eval","llm-evaluation","llm-evaluation-framework","llmops","pentesting","prompt-engineering","prompt-testing","prompts","rag","red-teaming","testing","vulnerability-scanners"],"archived":false,"github_pushed_at":"2026-08-01T23:47:56+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/promptfoo-promptfoo","markdown_url":"https://www.graphcanon.com/tools/promptfoo-promptfoo.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/promptfoo-promptfoo","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=promptfoo-promptfoo"}},{"type":"integrates_with","direction":"out","explanation":"Both tools focus on evaluating LLMs, lmms-eval for multimodal evaluations and promptfoo for evaluating and red-teaming LLM apps.","successor_context":null,"tool":{"slug":"promptfoo-promptfoo","name":"promptfoo","tagline":"Tool for evaluating prompts and AI agents by comparing performance across various models and red teaming.","github_url":"https://github.com/promptfoo/promptfoo","owner":"promptfoo","repo":"promptfoo","owner_avatar_url":"https://avatars.githubusercontent.com/u/137907881?v=4","primary_language":"TypeScript","stars":23838,"forks":2147,"topics":["ci","ci-cd","cicd","evaluation","evaluation-framework","llm","llm-eval","llm-evaluation","llm-evaluation-framework","llmops","pentesting","prompt-engineering","prompt-testing","prompts","rag","red-teaming","testing","vulnerability-scanners"],"archived":false,"github_pushed_at":"2026-08-01T23:47:56+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/promptfoo-promptfoo","markdown_url":"https://www.graphcanon.com/tools/promptfoo-promptfoo.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/promptfoo-promptfoo","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=promptfoo-promptfoo"}},{"type":"integrates_with","direction":"out","explanation":"Both tools aim at observability of AI systems, with lmms-eval focusing on multimodal evaluations while LMNR provides a broader platform for AI agent observability.","successor_context":null,"tool":{"slug":"lmnr-ai-lmnr","name":"lmnr","tagline":"Open-source observability platform for AI agents.","github_url":"https://github.com/lmnr-ai/lmnr","owner":"lmnr-ai","repo":"lmnr","owner_avatar_url":"https://avatars.githubusercontent.com/u/161496104?v=4","primary_language":"TypeScript","stars":3183,"forks":223,"topics":["agent-observability","agents","ai","ai-observability","aiops","analytics","developer-tools","evals","evaluation","llm-evaluation","llm-observability","llmops","monitoring","observability","open-source","rust","rust-lang","self-hosted","ts","typescript"],"archived":false,"github_pushed_at":"2026-08-20T09:30:48+00:00","maintenance_label":"Very active","stars_delta_30d":80,"url":"https://www.graphcanon.com/tools/lmnr-ai-lmnr","markdown_url":"https://www.graphcanon.com/tools/lmnr-ai-lmnr.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/lmnr-ai-lmnr","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=lmnr-ai-lmnr"}},{"type":"alternative","direction":"out","explanation":"Both tools focus on evaluating LLMs, but they offer different functionalities and approaches. Langfuse offers a broader platform for AI engineering with observability features, whereas lmms-eval is specialized in multimodal evaluations.","successor_context":null,"tool":{"slug":"langfuse-langfuse","name":"langfuse","tagline":"Open source AI engineering platform: LLM evals, observability, metrics, prompt management, playground, datasets","github_url":"https://github.com/langfuse/langfuse","owner":"langfuse","repo":"langfuse","owner_avatar_url":"https://avatars.githubusercontent.com/u/134601687?v=4","primary_language":"TypeScript","stars":32271,"forks":3466,"topics":["analytics","autogen","evaluation","langchain","large-language-models","llama-index","llm","llm-evaluation","llm-observability","llmops","monitoring","observability","open-source","openai","playground","prompt-engineering","prompt-management","self-hosted","ycombinator"],"archived":false,"github_pushed_at":"2026-07-31T22:58:07+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/langfuse-langfuse","markdown_url":"https://www.graphcanon.com/tools/langfuse-langfuse.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/langfuse-langfuse","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=langfuse-langfuse"}},{"type":"alternative","direction":"out","explanation":"Both tools offer observability solutions for ML and LLM models, but Evidently is an open-source framework tailored toward broader ML applications while lmms-eval focuses specifically on multimodal evaluation across various data types.","successor_context":null,"tool":{"slug":"evidentlyai-evidently","name":"evidently","tagline":"An open-source ML and LLM observability framework.","github_url":"https://github.com/evidentlyai/evidently","owner":"evidentlyai","repo":"evidently","owner_avatar_url":"https://avatars.githubusercontent.com/u/75031056?v=4","primary_language":"Jupyter Notebook","stars":7790,"forks":895,"topics":["data-drift","data-quality","data-science","data-validation","generative-ai","hacktoberfest","html-report","jupyter-notebook","llm","llmops","machine-learning","mlops","model-monitoring","pandas-dataframe"],"archived":false,"github_pushed_at":"2026-08-05T16:29:57+00:00","maintenance_label":"Very active","stars_delta_30d":117,"url":"https://www.graphcanon.com/tools/evidentlyai-evidently","markdown_url":"https://www.graphcanon.com/tools/evidentlyai-evidently.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/evidentlyai-evidently","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=evidentlyai-evidently"}},{"type":"related","direction":"in","explanation":null,"successor_context":null,"tool":{"slug":"prometheus-eval-prometheus-eval","name":"prometheus-eval","tagline":"Evaluate your LLM's response with Prometheus and GPT4","github_url":"https://github.com/prometheus-eval/prometheus-eval","owner":"prometheus-eval","repo":"prometheus-eval","owner_avatar_url":"https://avatars.githubusercontent.com/u/167460660?v=4","primary_language":"Python","stars":1107,"forks":68,"topics":["evaluation","gpt4","litellm","llm","llm-as-a-judge","llm-as-evaluator","llmops","python","vllm"],"archived":false,"github_pushed_at":"2025-04-25T03:58:37+00:00","maintenance_label":"Dormant","stars_delta_30d":5,"url":"https://www.graphcanon.com/tools/prometheus-eval-prometheus-eval","markdown_url":"https://www.graphcanon.com/tools/prometheus-eval-prometheus-eval.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/prometheus-eval-prometheus-eval","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=prometheus-eval-prometheus-eval"}},{"type":"successor","direction":"in","explanation":"VLMEvalKit has a 'successor' relationship to 'lmms-eval' because while lmms-eval provides a unified evaluation framework for multimodal large language models across various media types, VLMEvalKit specifically focuses on and simplifies the evaluation of vision-language models through dedicated generation-based methods and exact matching techniques, offering more specialized capabilities within a细分","successor_context":null,"tool":{"slug":"open-compass-vlmevalkit","name":"VLMEvalKit","tagline":"An open-source evaluation toolkit for large vision-language models","github_url":"https://github.com/open-compass/VLMEvalKit","owner":"open-compass","repo":"VLMEvalKit","owner_avatar_url":"https://avatars.githubusercontent.com/u/143521324?v=4","primary_language":"Python","stars":4345,"forks":745,"topics":["chatgpt","claude","clip","computer-vision","evaluation","gemini","gpt","gpt-4v","gpt4","large-language-models","llava","llm","multi-modal","openai","openai-api","pytorch","qwen","vit","vqa"],"archived":false,"github_pushed_at":"2026-08-17T17:16:26+00:00","maintenance_label":"Very active","stars_delta_30d":60,"url":"https://www.graphcanon.com/tools/open-compass-vlmevalkit","markdown_url":"https://www.graphcanon.com/tools/open-compass-vlmevalkit.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/open-compass-vlmevalkit","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=open-compass-vlmevalkit"}},{"type":"integrates_with","direction":"in","explanation":"The LMMS-evolution toolkit could be used for evaluating models highlighted in this repository.","successor_context":null,"tool":{"slug":"wangrongsheng-awesome-llm-resources","name":"awesome-LLM-resources","tagline":"Summary of the world's best LLM resources.","github_url":"https://github.com/WangRongsheng/awesome-LLM-resources","owner":"WangRongsheng","repo":"awesome-LLM-resources","owner_avatar_url":"https://avatars.githubusercontent.com/u/55651568?v=4","primary_language":null,"stars":8845,"forks":950,"topics":["awesome-list","book","course","large-language-models","llama","llm","mistral","openai","qwen","rag","retrieval-augmented-generation","webui"],"archived":false,"github_pushed_at":"2026-08-14T15:54:28+00:00","maintenance_label":"Very active","stars_delta_30d":142,"url":"https://www.graphcanon.com/tools/wangrongsheng-awesome-llm-resources","markdown_url":"https://www.graphcanon.com/tools/wangrongsheng-awesome-llm-resources.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/wangrongsheng-awesome-llm-resources","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=wangrongsheng-awesome-llm-resources"}},{"type":"alternative","direction":"in","explanation":"VLMEvalKit and lmms-eval both serve as evaluation toolkits for multimodal large language models, addressing the same problem with potentially different approaches or features.","successor_context":null,"tool":{"slug":"open-compass-vlmevalkit","name":"VLMEvalKit","tagline":"An open-source evaluation toolkit for large vision-language models","github_url":"https://github.com/open-compass/VLMEvalKit","owner":"open-compass","repo":"VLMEvalKit","owner_avatar_url":"https://avatars.githubusercontent.com/u/143521324?v=4","primary_language":"Python","stars":4345,"forks":745,"topics":["chatgpt","claude","clip","computer-vision","evaluation","gemini","gpt","gpt-4v","gpt4","large-language-models","llava","llm","multi-modal","openai","openai-api","pytorch","qwen","vit","vqa"],"archived":false,"github_pushed_at":"2026-08-17T17:16:26+00:00","maintenance_label":"Very active","stars_delta_30d":60,"url":"https://www.graphcanon.com/tools/open-compass-vlmevalkit","markdown_url":"https://www.graphcanon.com/tools/open-compass-vlmevalkit.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/open-compass-vlmevalkit","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=open-compass-vlmevalkit"}},{"type":"integrates_with","direction":"in","explanation":"NExT-GPT can be evaluated using lmms-eval, which is a unified evaluation toolkit for multimodal large language models.","successor_context":null,"tool":{"slug":"next-gpt-next-gpt","name":"NExT-GPT","tagline":"Code and models for ICML 2024 paper on multimodal large language model","github_url":"https://github.com/NExT-GPT/NExT-GPT","owner":"NExT-GPT","repo":"NExT-GPT","owner_avatar_url":"https://avatars.githubusercontent.com/u/143576855?v=4","primary_language":"Python","stars":3637,"forks":359,"topics":["chatgpt","foundation-models","gpt-4","instruction-tuning","large-language-models","llm","mllm","multi-modal-chatgpt","multimodal","visual-language-learning"],"archived":false,"github_pushed_at":"2025-05-13T09:57:47+00:00","maintenance_label":"Dormant","stars_delta_30d":-1,"url":"https://www.graphcanon.com/tools/next-gpt-next-gpt","markdown_url":"https://www.graphcanon.com/tools/next-gpt-next-gpt.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/next-gpt-next-gpt","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=next-gpt-next-gpt"}},{"type":"related","direction":"in","explanation":"Both deal with multimodal LLMs; MGM can be evaluated using lmms-eval toolkit, though they serve different functions.","successor_context":null,"tool":{"slug":"jia-lab-research-mgm","name":"MGM","tagline":"Mini-Gemini: Mining the Potential of Multi-modality Vision Language Models","github_url":"https://github.com/JIA-Lab-research/MGM","owner":"JIA-Lab-research","repo":"MGM","owner_avatar_url":"https://avatars.githubusercontent.com/u/64006090?v=4","primary_language":"Python","stars":3331,"forks":276,"topics":["generation","large-language-models","vision-language-model"],"archived":false,"github_pushed_at":"2024-05-04T14:36:51+00:00","maintenance_label":"Dormant","stars_delta_30d":1,"url":"https://www.graphcanon.com/tools/jia-lab-research-mgm","markdown_url":"https://www.graphcanon.com/tools/jia-lab-research-mgm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/jia-lab-research-mgm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=jia-lab-research-mgm"}}],"neighbours":[{"slug":"confident-ai-deepeval","name":"deepeval","tagline":"LLM Evaluation Framework.","github_url":"https://github.com/confident-ai/deepeval","owner":"confident-ai","repo":"deepeval","owner_avatar_url":"https://avatars.githubusercontent.com/u/130858411?v=4","primary_language":"Python","stars":17226,"forks":1736,"topics":["evaluation-framework","evaluation-metrics","llm-evaluation","llm-evaluation-framework","llm-evaluation-metrics","python"],"archived":false,"github_pushed_at":"2026-07-27T11:33:31+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/confident-ai-deepeval","markdown_url":"https://www.graphcanon.com/tools/confident-ai-deepeval.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/confident-ai-deepeval","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=confident-ai-deepeval","shared_categories":["evaluation-observability"]},{"slug":"eleutherai-lm-evaluation-harness","name":"lm-evaluation-harness","tagline":"A framework for few-shot evaluation of language models.","github_url":"https://github.com/EleutherAI/lm-evaluation-harness","owner":"EleutherAI","repo":"lm-evaluation-harness","owner_avatar_url":"https://avatars.githubusercontent.com/u/68924597?v=4","primary_language":"Python","stars":13560,"forks":3467,"topics":["evaluation-framework","language-model","transformer"],"archived":false,"github_pushed_at":"2026-07-13T20:18:15+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/eleutherai-lm-evaluation-harness","markdown_url":"https://www.graphcanon.com/tools/eleutherai-lm-evaluation-harness.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/eleutherai-lm-evaluation-harness","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=eleutherai-lm-evaluation-harness","shared_categories":["evaluation-observability"]},{"slug":"eugeneyan-open-llms","name":"open-llms","tagline":"A list of open LLMs available for commercial use.","github_url":"https://github.com/eugeneyan/open-llms","owner":"eugeneyan","repo":"open-llms","owner_avatar_url":"https://avatars.githubusercontent.com/u/6831355?v=4","primary_language":null,"stars":12849,"forks":985,"topics":["commercial","large-language-models","llm","llms"],"archived":false,"github_pushed_at":"2025-02-13T06:37:12+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/eugeneyan-open-llms","markdown_url":"https://www.graphcanon.com/tools/eugeneyan-open-llms.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/eugeneyan-open-llms","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=eugeneyan-open-llms","shared_categories":[]},{"slug":"steven2358-awesome-generative-ai","name":"awesome-generative-ai","tagline":"A curated list of modern Generative Artificial Intelligence projects and services","github_url":"https://github.com/steven2358/awesome-generative-ai","owner":"steven2358","repo":"awesome-generative-ai","owner_avatar_url":"https://avatars.githubusercontent.com/u/164072?v=4","primary_language":null,"stars":12501,"forks":1990,"topics":["ai","artificial-intelligence","awesome","awesome-list","generative-ai","generative-art","large-language-models","llm"],"archived":false,"github_pushed_at":"2026-08-03T10:58:05+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/steven2358-awesome-generative-ai","markdown_url":"https://www.graphcanon.com/tools/steven2358-awesome-generative-ai.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/steven2358-awesome-generative-ai","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=steven2358-awesome-generative-ai","shared_categories":[]},{"slug":"oumi-ai-oumi","name":"oumi","tagline":"Easily fine-tune, evaluate and deploy open source LLMs/VLMs","github_url":"https://github.com/oumi-ai/oumi","owner":"oumi-ai","repo":"oumi","owner_avatar_url":"https://avatars.githubusercontent.com/u/167452922?v=4","primary_language":"Python","stars":9359,"forks":780,"topics":["dpo","evaluation","fine-tuning","gpt-oss","gpt-oss-120b","gpt-oss-20b","inference","llama","llms","sft","slms","vlms"],"archived":false,"github_pushed_at":"2026-07-24T05:44:23+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/oumi-ai-oumi","markdown_url":"https://www.graphcanon.com/tools/oumi-ai-oumi.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/oumi-ai-oumi","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=oumi-ai-oumi","shared_categories":["evaluation-observability"]},{"slug":"wangrongsheng-awesome-llm-resources","name":"awesome-LLM-resources","tagline":"Summary of the world's best LLM resources.","github_url":"https://github.com/WangRongsheng/awesome-LLM-resources","owner":"WangRongsheng","repo":"awesome-LLM-resources","owner_avatar_url":"https://avatars.githubusercontent.com/u/55651568?v=4","primary_language":null,"stars":8845,"forks":950,"topics":["awesome-list","book","course","large-language-models","llama","llm","mistral","openai","qwen","rag","retrieval-augmented-generation","webui"],"archived":false,"github_pushed_at":"2026-08-14T15:54:28+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/wangrongsheng-awesome-llm-resources","markdown_url":"https://www.graphcanon.com/tools/wangrongsheng-awesome-llm-resources.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/wangrongsheng-awesome-llm-resources","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=wangrongsheng-awesome-llm-resources","shared_categories":["evaluation-observability"]},{"slug":"luban-agi-awesome-aigc-tutorials","name":"Awesome-AIGC-Tutorials","tagline":"Curated tutorials and resources for Large Language Models, AI Painting, and more","github_url":"https://github.com/luban-agi/Awesome-AIGC-Tutorials","owner":"luban-agi","repo":"Awesome-AIGC-Tutorials","owner_avatar_url":"https://avatars.githubusercontent.com/u/142385464?v=4","primary_language":null,"stars":4522,"forks":303,"topics":["ai","aigc","awesome","chatgpt","courses-resource","deep-learning","llm","midjourney","multimodal","nlp","prompt-engineering","stable-diffusion","tutorials"],"archived":false,"github_pushed_at":"2024-03-31T09:18:04+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/luban-agi-awesome-aigc-tutorials","markdown_url":"https://www.graphcanon.com/tools/luban-agi-awesome-aigc-tutorials.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/luban-agi-awesome-aigc-tutorials","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=luban-agi-awesome-aigc-tutorials","shared_categories":[]},{"slug":"nyldn-claude-octopus","name":"claude-octopus","tagline":"Surface AI blindspots before you ship","github_url":"https://github.com/nyldn/claude-octopus","owner":"nyldn","repo":"claude-octopus","owner_avatar_url":"https://avatars.githubusercontent.com/u/4805949?v=4","primary_language":"Shell","stars":3962,"forks":374,"topics":["ai-agents","ai-orchestration","claude-code","claude-code-plugin","codex","copilot","developer-tools","double-diamond","gemini","multi-ai","multi-llm","ollama"],"archived":false,"github_pushed_at":"2026-08-13T23:28:23+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/nyldn-claude-octopus","markdown_url":"https://www.graphcanon.com/tools/nyldn-claude-octopus.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/nyldn-claude-octopus","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=nyldn-claude-octopus","shared_categories":[]},{"slug":"hegelai-prompttools","name":"prompttools","tagline":"Open-source tools for prompt testing and experimentation","github_url":"https://github.com/hegelai/prompttools","owner":"hegelai","repo":"prompttools","owner_avatar_url":"https://avatars.githubusercontent.com/u/136523567?v=4","primary_language":"Python","stars":3046,"forks":255,"topics":["deep-learning","developer-tools","embeddings","large-language-models","llms","machine-learning","prompt-engineering","python","vector-search"],"archived":false,"github_pushed_at":"2026-02-11T03:24:04+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/hegelai-prompttools","markdown_url":"https://www.graphcanon.com/tools/hegelai-prompttools.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/hegelai-prompttools","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=hegelai-prompttools","shared_categories":[]},{"slug":"openinfer-project-openinfer","name":"openinfer","tagline":"Pure Rust CUDA LLM inference engine serving multiple models including Qwen3 and Kimi-K2","github_url":"https://github.com/openinfer-project/openinfer","owner":"openinfer-project","repo":"openinfer","owner_avatar_url":"https://avatars.githubusercontent.com/u/292134277?v=4","primary_language":"Rust","stars":585,"forks":89,"topics":["cuda","cuda-kernels","deepseek","gpu","inference","inference-engine","kimi","kimi-k2","kv-cache","llm","llm-inference","llm-serving","model-serving","moe","openai-api","paged-attention","qwen","qwen3","rust","vllm"],"archived":false,"github_pushed_at":"2026-07-25T14:08:34+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/openinfer-project-openinfer","markdown_url":"https://www.graphcanon.com/tools/openinfer-project-openinfer.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/openinfer-project-openinfer","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=openinfer-project-openinfer","shared_categories":[]},{"slug":"declare-lab-instruct-eval","name":"instruct-eval","tagline":"Quantitative evaluation for instruction-tuned language models","github_url":"https://github.com/declare-lab/instruct-eval","owner":"declare-lab","repo":"instruct-eval","owner_avatar_url":"https://avatars.githubusercontent.com/u/59164695?v=4","primary_language":"Python","stars":552,"forks":45,"topics":["instruct-tuning","llm"],"archived":false,"github_pushed_at":"2024-03-10T05:00:00+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/declare-lab-instruct-eval","markdown_url":"https://www.graphcanon.com/tools/declare-lab-instruct-eval.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/declare-lab-instruct-eval","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=declare-lab-instruct-eval","shared_categories":["evaluation-observability"]},{"slug":"curated-awesome-lists-awesome-llms-fine-tuning","name":"awesome-llms-fine-tuning","tagline":"A comprehensive collection of resources for fine-tuning Large Language Models.","github_url":"https://github.com/Curated-Awesome-Lists/awesome-llms-fine-tuning","owner":"Curated-Awesome-Lists","repo":"awesome-llms-fine-tuning","owner_avatar_url":"https://avatars.githubusercontent.com/u/142611331?v=4","primary_language":null,"stars":525,"forks":78,"topics":["ai","awesome-list","deep-learning","fine-tuning","gpt","large-language-models","llms","machine-learning","nlp","transformers"],"archived":false,"github_pushed_at":"2024-12-02T20:11:14+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/curated-awesome-lists-awesome-llms-fine-tuning","markdown_url":"https://www.graphcanon.com/tools/curated-awesome-lists-awesome-llms-fine-tuning.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/curated-awesome-lists-awesome-llms-fine-tuning","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=curated-awesome-lists-awesome-llms-fine-tuning","shared_categories":[]}]}}