{"data":{"node":{"slug":"uptrain-ai-uptrain","name":"uptrain","tagline":"Unified platform for evaluating and improving Generative AI applications","github_url":"https://github.com/uptrain-ai/uptrain","owner":"uptrain-ai","repo":"uptrain","owner_avatar_url":"https://avatars.githubusercontent.com/u/114582870?v=4","primary_language":"Python","stars":2359,"forks":204,"topics":["autoevaluation","evaluation","experimentation","hallucination-detection","jailbreak-detection","llm-eval","llm-prompting","llm-test","llmops","machine-learning","monitoring","openai-evals","prompt-engineering","root-cause-analysis"],"archived":false,"github_pushed_at":"2024-08-18T13:30:44+00:00","maintenance_label":"Dormant","stars_delta_30d":4,"url":"https://www.graphcanon.com/tools/uptrain-ai-uptrain","markdown_url":"https://www.graphcanon.com/tools/uptrain-ai-uptrain.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/uptrain-ai-uptrain","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=uptrain-ai-uptrain"},"categories":[{"slug":"evaluation-observability","name":"Evaluation & Observability","url":"https://www.graphcanon.com/categories/evaluation-observability","markdown_url":"https://www.graphcanon.com/categories/evaluation-observability.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/evaluation-observability"}],"tags":[{"slug":"autoevaluation","name":"autoevaluation"},{"slug":"evaluation","name":"evaluation"},{"slug":"experimentation","name":"experimentation"},{"slug":"hallucination-detection","name":"hallucination-detection"},{"slug":"jailbreak-detection","name":"jailbreak-detection"},{"slug":"llm-eval","name":"llm-eval"},{"slug":"llm-prompting","name":"llm-prompting"},{"slug":"llm-test","name":"llm-test"}],"edges":[{"type":"related","direction":"out","explanation":null,"successor_context":null,"tool":{"slug":"promptfoo-promptfoo","name":"promptfoo","tagline":"Tool for evaluating prompts and AI agents by comparing performance across various models and red teaming.","github_url":"https://github.com/promptfoo/promptfoo","owner":"promptfoo","repo":"promptfoo","owner_avatar_url":"https://avatars.githubusercontent.com/u/137907881?v=4","primary_language":"TypeScript","stars":23838,"forks":2147,"topics":["ci","ci-cd","cicd","evaluation","evaluation-framework","llm","llm-eval","llm-evaluation","llm-evaluation-framework","llmops","pentesting","prompt-engineering","prompt-testing","prompts","rag","red-teaming","testing","vulnerability-scanners"],"archived":false,"github_pushed_at":"2026-08-01T23:47:56+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/promptfoo-promptfoo","markdown_url":"https://www.graphcanon.com/tools/promptfoo-promptfoo.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/promptfoo-promptfoo","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=promptfoo-promptfoo"}},{"type":"alternative","direction":"out","explanation":"Both UpTrain and Langfuse focus on evaluating and improving LLM applications but offer different functionalities and approaches.","successor_context":null,"tool":{"slug":"langfuse-langfuse","name":"langfuse","tagline":"Open source AI engineering platform: LLM evals, observability, metrics, prompt management, playground, datasets","github_url":"https://github.com/langfuse/langfuse","owner":"langfuse","repo":"langfuse","owner_avatar_url":"https://avatars.githubusercontent.com/u/134601687?v=4","primary_language":"TypeScript","stars":32271,"forks":3466,"topics":["analytics","autogen","evaluation","langchain","large-language-models","llama-index","llm","llm-evaluation","llm-observability","llmops","monitoring","observability","open-source","openai","playground","prompt-engineering","prompt-management","self-hosted","ycombinator"],"archived":false,"github_pushed_at":"2026-07-31T22:58:07+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/langfuse-langfuse","markdown_url":"https://www.graphcanon.com/tools/langfuse-langfuse.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/langfuse-langfuse","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=langfuse-langfuse"}},{"type":"alternative","direction":"out","explanation":"Evidently also provides observability and evaluation features targeting ML and LLM models, akin to UpTrain's scope of operations.","successor_context":null,"tool":{"slug":"evidentlyai-evidently","name":"evidently","tagline":"An open-source ML and LLM observability framework.","github_url":"https://github.com/evidentlyai/evidently","owner":"evidentlyai","repo":"evidently","owner_avatar_url":"https://avatars.githubusercontent.com/u/75031056?v=4","primary_language":"Jupyter Notebook","stars":7790,"forks":895,"topics":["data-drift","data-quality","data-science","data-validation","generative-ai","hacktoberfest","html-report","jupyter-notebook","llm","llmops","machine-learning","mlops","model-monitoring","pandas-dataframe"],"archived":false,"github_pushed_at":"2026-08-05T16:29:57+00:00","maintenance_label":"Very active","stars_delta_30d":117,"url":"https://www.graphcanon.com/tools/evidentlyai-evidently","markdown_url":"https://www.graphcanon.com/tools/evidentlyai-evidently.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/evidentlyai-evidently","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=evidentlyai-evidently"}},{"type":"alternative","direction":"out","explanation":"Phoenix from Arize AI is also centered around the observability and evaluation of ML models, making it comparable in purpose to UpTrain while differing in specific features.","successor_context":null,"tool":{"slug":"arize-ai-phoenix","name":"phoenix","tagline":"AI Observability & Evaluation","github_url":"https://github.com/Arize-ai/phoenix","owner":"Arize-ai","repo":"phoenix","owner_avatar_url":"https://avatars.githubusercontent.com/u/59858760?v=4","primary_language":"Python","stars":10847,"forks":1028,"topics":["agents","ai-monitoring","ai-observability","aiengineering","anthropic","datasets","evals","langchain","llamaindex","llm-eval","llm-evaluation","llmops","llms","openai","prompt-engineering","smolagents"],"archived":false,"github_pushed_at":"2026-08-01T11:48:49+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/arize-ai-phoenix","markdown_url":"https://www.graphcanon.com/tools/arize-ai-phoenix.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/arize-ai-phoenix","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=arize-ai-phoenix"}},{"type":"related","direction":"out","explanation":null,"successor_context":null,"tool":{"slug":"comet-ml-opik","name":"opik","tagline":"Debug, evaluate, and monitor your LLM applications with comprehensive tracing and production-ready dashboards","github_url":"https://github.com/comet-ml/opik","owner":"comet-ml","repo":"opik","owner_avatar_url":"https://avatars.githubusercontent.com/u/31487821?v=4","primary_language":"Python","stars":21177,"forks":1681,"topics":["evaluation","hacktoberfest","hacktoberfest2025","langchain","llama-index","llm","llm-evaluation","llm-observability","llmops","open-source","openai","playground","prompt-engineering"],"archived":false,"github_pushed_at":"2026-08-07T11:46:34+00:00","maintenance_label":"Very active","stars_delta_30d":767,"url":"https://www.graphcanon.com/tools/comet-ml-opik","markdown_url":"https://www.graphcanon.com/tools/comet-ml-opik.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/comet-ml-opik","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=comet-ml-opik"}},{"type":"related","direction":"out","explanation":null,"successor_context":null,"tool":{"slug":"agenta-ai-agenta","name":"agenta","tagline":"The open-source LLMOps platform for prompt management, evaluation, and observability.","github_url":"https://github.com/Agenta-AI/agenta","owner":"Agenta-AI","repo":"agenta","owner_avatar_url":"https://avatars.githubusercontent.com/u/127993667?v=4","primary_language":"TypeScript","stars":4445,"forks":609,"topics":["agent-builder","agent-observability","agent-orchestration","agent-workspace","agentic-ai","ai-agent","ai-agents","ai-automation","ai-skills-manager","ai-workflow-builder","harness","mcp","open-source","self-hosted","workflow-automation"],"archived":false,"github_pushed_at":"2026-08-07T10:41:36+00:00","maintenance_label":"Very active","stars_delta_30d":170,"url":"https://www.graphcanon.com/tools/agenta-ai-agenta","markdown_url":"https://www.graphcanon.com/tools/agenta-ai-agenta.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/agenta-ai-agenta","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=agenta-ai-agenta"}},{"type":"successor","direction":"out","explanation":"Ragas provides domain-specific evaluations and optimization tools specifically for Large Language Models (LLMs), while UpTrain offers a broader platform for both evaluating and improving generative AI systems including LLMs. The successor relationship from Ragas to UpTrain suggests an expansion in functionality from specialized LLM evaluation to a more comprehensive solution for various types of G","successor_context":null,"tool":{"slug":"vibrantlabsai-ragas","name":"ragas","tagline":"Supercharge Your LLM Application Evaluations 🚀","github_url":"https://github.com/vibrantlabsai/ragas","owner":"vibrantlabsai","repo":"ragas","owner_avatar_url":"https://avatars.githubusercontent.com/u/122604797?v=4","primary_language":"Python","stars":15388,"forks":1637,"topics":["evaluation","llm","llmops"],"archived":false,"github_pushed_at":"2026-02-24T07:47:19+00:00","maintenance_label":"Slowing","stars_delta_30d":470,"url":"https://www.graphcanon.com/tools/vibrantlabsai-ragas","markdown_url":"https://www.graphcanon.com/tools/vibrantlabsai-ragas.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/vibrantlabsai-ragas","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=vibrantlabsai-ragas"}},{"type":"related","direction":"out","explanation":null,"successor_context":null,"tool":{"slug":"traceloop-openllmetry","name":"openllmetry","tagline":"Open-source observability for GenAI and LLM applications based on OpenTelemetry.","github_url":"https://github.com/traceloop/openllmetry","owner":"traceloop","repo":"openllmetry","owner_avatar_url":"https://avatars.githubusercontent.com/u/125419530?v=4","primary_language":"Python","stars":7377,"forks":1047,"topics":["artifical-intelligence","datascience","generative-ai","good-first-issue","good-first-issues","help-wanted","llm","llmops","metrics","ml","model-monitoring","monitoring","observability","open-source","open-telemetry","opentelemetry","opentelemetry-python","python"],"archived":false,"github_pushed_at":"2026-08-10T08:49:01+00:00","maintenance_label":"Very active","stars_delta_30d":75,"url":"https://www.graphcanon.com/tools/traceloop-openllmetry","markdown_url":"https://www.graphcanon.com/tools/traceloop-openllmetry.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/traceloop-openllmetry","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=traceloop-openllmetry"}}],"neighbours":[{"slug":"lm-sys-fastchat","name":"FastChat","tagline":"An open platform for training, serving, and evaluating large language models","github_url":"https://github.com/lm-sys/FastChat","owner":"lm-sys","repo":"FastChat","owner_avatar_url":"https://avatars.githubusercontent.com/u/126381704?v=4","primary_language":"Python","stars":39517,"forks":4788,"topics":[],"archived":false,"github_pushed_at":"2026-05-01T00:25:53+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/lm-sys-fastchat","markdown_url":"https://www.graphcanon.com/tools/lm-sys-fastchat.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/lm-sys-fastchat","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=lm-sys-fastchat","shared_categories":["evaluation-observability"]},{"slug":"patchy631-ai-engineering-hub","name":"ai-engineering-hub","tagline":"Tutorials on LLMs, RAGs, and real-world AI agent applications","github_url":"https://github.com/patchy631/ai-engineering-hub","owner":"patchy631","repo":"ai-engineering-hub","owner_avatar_url":"https://avatars.githubusercontent.com/u/38653995?v=4","primary_language":"Jupyter Notebook","stars":37020,"forks":6107,"topics":["agents","ai","llms","machine-learning","mcp","rag"],"archived":false,"github_pushed_at":"2026-07-27T18:43:06+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/patchy631-ai-engineering-hub","markdown_url":"https://www.graphcanon.com/tools/patchy631-ai-engineering-hub.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/patchy631-ai-engineering-hub","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=patchy631-ai-engineering-hub","shared_categories":[]},{"slug":"linshenkx-prompt-optimizer","name":"prompt-optimizer","tagline":"An AI prompt optimizer for writing better prompts and getting better AI results.","github_url":"https://github.com/linshenkx/prompt-optimizer","owner":"linshenkx","repo":"prompt-optimizer","owner_avatar_url":"https://avatars.githubusercontent.com/u/32978552?v=4","primary_language":"TypeScript","stars":33144,"forks":3890,"topics":["ai-prompts","ai-tools","llm","prompt","prompt-engineering","prompt-optimization","prompt-optimizer","prompt-testing","prompt-toolkit","prompt-tuning"],"archived":false,"github_pushed_at":"2026-08-13T14:23:12+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/linshenkx-prompt-optimizer","markdown_url":"https://www.graphcanon.com/tools/linshenkx-prompt-optimizer.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/linshenkx-prompt-optimizer","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=linshenkx-prompt-optimizer","shared_categories":["evaluation-observability"]},{"slug":"promptfoo-promptfoo","name":"promptfoo","tagline":"Tool for evaluating prompts and AI agents by comparing performance across various models and red teaming.","github_url":"https://github.com/promptfoo/promptfoo","owner":"promptfoo","repo":"promptfoo","owner_avatar_url":"https://avatars.githubusercontent.com/u/137907881?v=4","primary_language":"TypeScript","stars":23838,"forks":2147,"topics":["ci","ci-cd","cicd","evaluation","evaluation-framework","llm","llm-eval","llm-evaluation","llm-evaluation-framework","llmops","pentesting","prompt-engineering","prompt-testing","prompts","rag","red-teaming","testing","vulnerability-scanners"],"archived":false,"github_pushed_at":"2026-08-01T23:47:56+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/promptfoo-promptfoo","markdown_url":"https://www.graphcanon.com/tools/promptfoo-promptfoo.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/promptfoo-promptfoo","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=promptfoo-promptfoo","shared_categories":["evaluation-observability"]},{"slug":"comet-ml-opik","name":"opik","tagline":"Debug, evaluate, and monitor your LLM applications with comprehensive tracing and production-ready dashboards","github_url":"https://github.com/comet-ml/opik","owner":"comet-ml","repo":"opik","owner_avatar_url":"https://avatars.githubusercontent.com/u/31487821?v=4","primary_language":"Python","stars":21177,"forks":1681,"topics":["evaluation","hacktoberfest","hacktoberfest2025","langchain","llama-index","llm","llm-evaluation","llm-observability","llmops","open-source","openai","playground","prompt-engineering"],"archived":false,"github_pushed_at":"2026-08-07T11:46:34+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/comet-ml-opik","markdown_url":"https://www.graphcanon.com/tools/comet-ml-opik.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/comet-ml-opik","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=comet-ml-opik","shared_categories":["evaluation-observability"]},{"slug":"confident-ai-deepeval","name":"deepeval","tagline":"LLM Evaluation Framework.","github_url":"https://github.com/confident-ai/deepeval","owner":"confident-ai","repo":"deepeval","owner_avatar_url":"https://avatars.githubusercontent.com/u/130858411?v=4","primary_language":"Python","stars":17226,"forks":1736,"topics":["evaluation-framework","evaluation-metrics","llm-evaluation","llm-evaluation-framework","llm-evaluation-metrics","python"],"archived":false,"github_pushed_at":"2026-07-27T11:33:31+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/confident-ai-deepeval","markdown_url":"https://www.graphcanon.com/tools/confident-ai-deepeval.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/confident-ai-deepeval","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=confident-ai-deepeval","shared_categories":["evaluation-observability"]},{"slug":"raga-ai-hub-ragaai-catalyst","name":"RagaAI-Catalyst","tagline":"Python SDK for AI agent observability and evaluation","github_url":"https://github.com/raga-ai-hub/RagaAI-Catalyst","owner":"raga-ai-hub","repo":"RagaAI-Catalyst","owner_avatar_url":"https://avatars.githubusercontent.com/u/161833182?v=4","primary_language":"Python","stars":16148,"forks":3565,"topics":["agentic-ai","agentic-ai-development","agentneo","agents","ai-agent-monitoring","ai-application-debugging","ai-evaluation-tools","ai-performance-optimization","ai-tool-interaction-monitoring","llm-testing","llm-tracing","llmops"],"archived":false,"github_pushed_at":"2026-02-11T14:43:33+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/raga-ai-hub-ragaai-catalyst","markdown_url":"https://www.graphcanon.com/tools/raga-ai-hub-ragaai-catalyst.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/raga-ai-hub-ragaai-catalyst","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=raga-ai-hub-ragaai-catalyst","shared_categories":["evaluation-observability"]},{"slug":"steven2358-awesome-generative-ai","name":"awesome-generative-ai","tagline":"A curated list of modern Generative Artificial Intelligence projects and services","github_url":"https://github.com/steven2358/awesome-generative-ai","owner":"steven2358","repo":"awesome-generative-ai","owner_avatar_url":"https://avatars.githubusercontent.com/u/164072?v=4","primary_language":null,"stars":12501,"forks":1990,"topics":["ai","artificial-intelligence","awesome","awesome-list","generative-ai","generative-art","large-language-models","llm"],"archived":false,"github_pushed_at":"2026-08-03T10:58:05+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/steven2358-awesome-generative-ai","markdown_url":"https://www.graphcanon.com/tools/steven2358-awesome-generative-ai.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/steven2358-awesome-generative-ai","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=steven2358-awesome-generative-ai","shared_categories":[]},{"slug":"kiln-ai-kiln","name":"Kiln","tagline":"Build, Evaluate, and Optimize AI Systems","github_url":"https://github.com/Kiln-AI/Kiln","owner":"Kiln-AI","repo":"Kiln","owner_avatar_url":"https://avatars.githubusercontent.com/u/178670964?v=4","primary_language":"Python","stars":4971,"forks":374,"topics":["ai","chain-of-thought","collaboration","dataset-generation","evals","evaluation","evaluation-framework","fine-tuning","machine-learning","macos","mcp","ml","ollama","openai","prompt","prompt-engineering","python","rlhf","synthetic-data","windows"],"archived":false,"github_pushed_at":"2026-07-23T07:36:00+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/kiln-ai-kiln","markdown_url":"https://www.graphcanon.com/tools/kiln-ai-kiln.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/kiln-ai-kiln","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=kiln-ai-kiln","shared_categories":["evaluation-observability"]},{"slug":"trigaten-learn-prompting","name":"Learn_Prompting","tagline":"Your Go-To Resource for Mastering Generative AI","github_url":"https://github.com/trigaten/Learn_Prompting","owner":"trigaten","repo":"Learn_Prompting","owner_avatar_url":"https://avatars.githubusercontent.com/u/42948753?v=4","primary_language":"MDX","stars":4726,"forks":668,"topics":["chatgpt","chatgpt-api","deep-learning","gpt-3","gpt-4","gpt-4-api","gpt3","large-language-models","llm","machine-learning","nlp","openai-api","prompt-engineering","prompt-toolkit","prompt-tuning","prompting","transformers"],"archived":false,"github_pushed_at":"2025-01-14T15:04:07+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/trigaten-learn-prompting","markdown_url":"https://www.graphcanon.com/tools/trigaten-learn-prompting.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/trigaten-learn-prompting","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=trigaten-learn-prompting","shared_categories":[]},{"slug":"nyldn-claude-octopus","name":"claude-octopus","tagline":"Surface AI blindspots before you ship","github_url":"https://github.com/nyldn/claude-octopus","owner":"nyldn","repo":"claude-octopus","owner_avatar_url":"https://avatars.githubusercontent.com/u/4805949?v=4","primary_language":"Shell","stars":3962,"forks":374,"topics":["ai-agents","ai-orchestration","claude-code","claude-code-plugin","codex","copilot","developer-tools","double-diamond","gemini","multi-ai","multi-llm","ollama"],"archived":false,"github_pushed_at":"2026-08-13T23:28:23+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/nyldn-claude-octopus","markdown_url":"https://www.graphcanon.com/tools/nyldn-claude-octopus.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/nyldn-claude-octopus","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=nyldn-claude-octopus","shared_categories":[]},{"slug":"nerogar-onetrainer","name":"OneTrainer","tagline":"A comprehensive tool for Diffusion model training","github_url":"https://github.com/Nerogar/OneTrainer","owner":"Nerogar","repo":"OneTrainer","owner_avatar_url":"https://avatars.githubusercontent.com/u/3390934?v=4","primary_language":"Python","stars":3126,"forks":315,"topics":["fine-tuning","image-model-training","lora","training"],"archived":false,"github_pushed_at":"2026-07-20T19:42:24+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/nerogar-onetrainer","markdown_url":"https://www.graphcanon.com/tools/nerogar-onetrainer.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/nerogar-onetrainer","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=nerogar-onetrainer","shared_categories":[]}]}}