{"data":{"slug":"confident-ai-deepeval","name":"deepeval","tagline":"LLM Evaluation Framework.","github_url":"https://github.com/confident-ai/deepeval","owner":"confident-ai","repo":"deepeval","owner_avatar_url":"https://avatars.githubusercontent.com/u/130858411?v=4","primary_language":"Python","stars":17226,"forks":1736,"topics":["evaluation-framework","evaluation-metrics","llm-evaluation","llm-evaluation-framework","llm-evaluation-metrics","python"],"archived":false,"github_pushed_at":"2026-07-27T11:33:31+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/confident-ai-deepeval","markdown_url":"https://www.graphcanon.com/tools/confident-ai-deepeval.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/confident-ai-deepeval","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=confident-ai-deepeval","description":"The LLM Evaluation Framework","homepage_url":"https://deepeval.com","license":"Apache-2.0","open_issues":404,"watchers":66,"ai_summary":"A Python-based framework for evaluating large language models with various metrics and evaluation methodologies.","readme_excerpt":"## Installation\n\nDeepeval works with **Python>=3.9+**.\n\n```\npip install -U deepeval\n```\n\n---\n\n# License\n\nDeepEval is licensed under Apache 2.0 - see the [LICENSE.md](https://github.com/confident-ai/deepeval/blob/main/LICENSE.md) file for details.","github_created_at":"2023-08-10T05:35:04+00:00","created_at":"2026-07-11T11:59:08.209204+00:00","updated_at":"2026-07-28T12:00:25.329211+00:00","categories":[{"slug":"evaluation-observability","name":"Evaluation & Observability","url":"https://www.graphcanon.com/categories/evaluation-observability","markdown_url":"https://www.graphcanon.com/categories/evaluation-observability.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/evaluation-observability"}],"tags":[{"slug":"evaluation","name":"evaluation"},{"slug":"llm-evaluation","name":"llm-evaluation"},{"slug":"metrics","name":"metrics"}],"trust":{"provenance":{"is_fork":false,"github_id":676829188,"owner_type":"Organization","methodology":"github_public_v1","parent_repo":null,"near_duplicate_slugs":[]},"computed_at":"2026-07-28T12:00:24.387Z","maintenance":{"label":"Very active","score":96,"methodology":"github_public_v1","releases_90d":4,"days_since_push":1,"last_release_at":"2026-07-12T09:31:22Z"},"security_summary":{"status":"no_lockfile","scanner":null,"low_count":0,"high_count":0,"last_scan_at":"2026-07-11T11:59:09.396Z","medium_count":0,"scan_profile":"none","critical_count":0}},"capability_facts":{"scan":{"source":"repo_scan","observed_at":"2026-07-28T12:00:25.032Z"},"has_cli":{"value":true,"source":"pyproject.toml:[project.scripts]","observed_at":"2026-07-28T12:00:25.032Z"},"languages":{"value":["python"],"source":"github.language+pyproject.toml","observed_at":"2026-07-28T12:00:25.032Z"},"license_spdx":{"value":"Apache-2.0","source":"github.license","observed_at":"2026-07-28T12:00:25.032Z"}},"decision_facts":{"hosting":null,"pricing":null,"requirements":{"notes":["Requires Python environment and familiarity with large language models to effectively utilize Deepeval's capabilities."]},"constraints":null,"when_to_use":["When developing large language models and you need a comprehensive evaluation framework to measure their performance across various metrics.","If your project requires detailed, nuanced analysis provided by Deepeval's specific evaluation methodologies tailored to large language models."],"when_not_to_use":["For small-scale applications that do not require the depth of metrics and evaluations offered by Deepeval, as it might be overkill.","In situations where there is a need for real-time performance monitoring, since Deepeval focuses more on post-development evaluation rather than continuous runtime analysis."],"source":"enrich:decision_facts","observed_at":"2026-07-17T13:06:53.954Z"},"constraint_facets":null,"decision_summary":[{"label":"Requirements","value":"Requires Python environment and familiarity with large language models to effectively utilize Deepeval's capabilities."},{"label":"Adopt for","value":"Deepeval is a Python-based framework designed for evaluating large language models with an array of metrics and evaluation methodologies."},{"label":"License detail","value":"Apache-2.0 License"}]}}