{"data":{"node":{"slug":"babelscape-alert","name":"ALERT","tagline":"A Comprehensive Benchmark for Assessing Large Language Models' Safety Through Red Teaming","github_url":"https://github.com/Babelscape/ALERT","owner":"Babelscape","repo":"ALERT","owner_avatar_url":"https://avatars.githubusercontent.com/u/90899893?v=4","primary_language":"Python","stars":59,"forks":8,"topics":["ai","artificial-intelligence","benchmark","bias-detection","llm","llm-evaluation","llm-safety","llm-safety-benchmark","nlp","nlp-machine-learning","red-teaming","safety-monitoring","transformers-models"],"archived":false,"github_pushed_at":"2024-09-20T08:29:57+00:00","maintenance_label":"Dormant","stars_delta_30d":0,"url":"https://www.graphcanon.com/tools/babelscape-alert","markdown_url":"https://www.graphcanon.com/tools/babelscape-alert.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/babelscape-alert","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=babelscape-alert"},"categories":[{"slug":"evaluation-observability","name":"Evaluation & Observability","url":"https://www.graphcanon.com/categories/evaluation-observability","markdown_url":"https://www.graphcanon.com/categories/evaluation-observability.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/evaluation-observability"}],"tags":[{"slug":"ai","name":"ai"},{"slug":"artificial-intelligence","name":"artificial-intelligence"},{"slug":"benchmark","name":"benchmark"},{"slug":"bias-detection","name":"bias-detection"},{"slug":"llm-evaluation","name":"llm-evaluation"},{"slug":"llm-safety","name":"llm-safety"},{"slug":"nlp","name":"nlp"},{"slug":"transformers-models","name":"transformers-models"}],"edges":[],"neighbours":[{"slug":"llm-attacks-llm-attacks","name":"llm-attacks","tagline":"Universal and Transferable Attacks on Aligned Language Models","github_url":"https://github.com/llm-attacks/llm-attacks","owner":"llm-attacks","repo":"llm-attacks","owner_avatar_url":"https://avatars.githubusercontent.com/u/140664770?v=4","primary_language":"Python","stars":4780,"forks":635,"topics":[],"archived":false,"github_pushed_at":"2024-08-02T06:02:18+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/llm-attacks-llm-attacks","markdown_url":"https://www.graphcanon.com/tools/llm-attacks-llm-attacks.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/llm-attacks-llm-attacks","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=llm-attacks-llm-attacks","shared_categories":["evaluation-observability"]},{"slug":"meta-llama-purplellama","name":"PurpleLlama","tagline":"Set of tools to assess and improve LLM security","github_url":"https://github.com/meta-llama/PurpleLlama","owner":"meta-llama","repo":"PurpleLlama","owner_avatar_url":"https://avatars.githubusercontent.com/u/153379578?v=4","primary_language":"Python","stars":4380,"forks":769,"topics":[],"archived":false,"github_pushed_at":"2026-08-18T01:15:09+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/meta-llama-purplellama","markdown_url":"https://www.graphcanon.com/tools/meta-llama-purplellama.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/meta-llama-purplellama","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=meta-llama-purplellama","shared_categories":["evaluation-observability"]},{"slug":"corca-ai-awesome-llm-security","name":"awesome-llm-security","tagline":"A curation of tools, documents and projects about LLM Security","github_url":"https://github.com/corca-ai/awesome-llm-security","owner":"corca-ai","repo":"awesome-llm-security","owner_avatar_url":"https://avatars.githubusercontent.com/u/72978860?v=4","primary_language":null,"stars":1692,"forks":347,"topics":["awesome","awesome-list","llm","security"],"archived":false,"github_pushed_at":"2025-08-20T01:27:47+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/corca-ai-awesome-llm-security","markdown_url":"https://www.graphcanon.com/tools/corca-ai-awesome-llm-security.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/corca-ai-awesome-llm-security","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=corca-ai-awesome-llm-security","shared_categories":["evaluation-observability"]},{"slug":"livecodebench-livecodebench","name":"LiveCodeBench","tagline":"Holistic and contamination-free evaluation of large language models for code","github_url":"https://github.com/LiveCodeBench/LiveCodeBench","owner":"LiveCodeBench","repo":"LiveCodeBench","owner_avatar_url":"https://avatars.githubusercontent.com/u/161278213?v=4","primary_language":"Python","stars":936,"forks":199,"topics":["code-execution","code-generation","code-llms","code-repair","gpt-4","test-generation"],"archived":false,"github_pushed_at":"2025-07-16T00:58:38+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/livecodebench-livecodebench","markdown_url":"https://www.graphcanon.com/tools/livecodebench-livecodebench.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/livecodebench-livecodebench","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=livecodebench-livecodebench","shared_categories":["evaluation-observability"]},{"slug":"jailbreakbench-jailbreakbench","name":"jailbreakbench","tagline":"An Open Robustness Benchmark for Jailbreaking Language Models","github_url":"https://github.com/JailbreakBench/jailbreakbench","owner":"JailbreakBench","repo":"jailbreakbench","owner_avatar_url":"https://avatars.githubusercontent.com/u/151564101?v=4","primary_language":"Python","stars":665,"forks":78,"topics":[],"archived":false,"github_pushed_at":"2025-04-04T11:30:46+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/jailbreakbench-jailbreakbench","markdown_url":"https://www.graphcanon.com/tools/jailbreakbench-jailbreakbench.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/jailbreakbench-jailbreakbench","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=jailbreakbench-jailbreakbench","shared_categories":["evaluation-observability"]},{"slug":"sherdencooper-gptfuzz","name":"GPTFuzz","tagline":"Red Teaming Large Language Models with Auto-Generated Jailbreak Prompts","github_url":"https://github.com/sherdencooper/GPTFuzz","owner":"sherdencooper","repo":"GPTFuzz","owner_avatar_url":"https://avatars.githubusercontent.com/u/37368657?v=4","primary_language":"Python","stars":608,"forks":87,"topics":[],"archived":false,"github_pushed_at":"2026-02-27T17:19:03+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/sherdencooper-gptfuzz","markdown_url":"https://www.graphcanon.com/tools/sherdencooper-gptfuzz.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/sherdencooper-gptfuzz","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=sherdencooper-gptfuzz","shared_categories":["evaluation-observability"]},{"slug":"deadbits-vigil-llm","name":"vigil-llm","tagline":"Detect prompt injections and other risky inputs in LLMs","github_url":"https://github.com/deadbits/vigil-llm","owner":"deadbits","repo":"vigil-llm","owner_avatar_url":"https://avatars.githubusercontent.com/u/1332757?v=4","primary_language":"Python","stars":496,"forks":56,"topics":["adversarial-attacks","adversarial-machine-learning","large-language-models","llm-security","llmops","prompt-injection","security-tools","yara-scanner"],"archived":false,"github_pushed_at":"2024-01-31T18:43:41+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/deadbits-vigil-llm","markdown_url":"https://www.graphcanon.com/tools/deadbits-vigil-llm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/deadbits-vigil-llm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=deadbits-vigil-llm","shared_categories":["evaluation-observability"]},{"slug":"liu00222-open-prompt-injection","name":"Open-Prompt-Injection","tagline":"Benchmark and toolkit for prompt injection attacks and defenses in LLMs","github_url":"https://github.com/liu00222/Open-Prompt-Injection","owner":"liu00222","repo":"Open-Prompt-Injection","owner_avatar_url":"https://avatars.githubusercontent.com/u/42081599?v=4","primary_language":"Python","stars":489,"forks":78,"topics":["llm","llm-security","llms","prompt-injection","prompt-injection-tool","security-and-privacy"],"archived":false,"github_pushed_at":"2025-10-29T17:11:34+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/liu00222-open-prompt-injection","markdown_url":"https://www.graphcanon.com/tools/liu00222-open-prompt-injection.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/liu00222-open-prompt-injection","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=liu00222-open-prompt-injection","shared_categories":["evaluation-observability"]},{"slug":"llm-tuning-safety-llms-finetuning-safety","name":"LLMs-Finetuning-Safety","tagline":"Demonstrates safety risks in fine-tuning GPT-3.5 Turbo with adversarial examples","github_url":"https://github.com/LLM-Tuning-Safety/LLMs-Finetuning-Safety","owner":"LLM-Tuning-Safety","repo":"LLMs-Finetuning-Safety","owner_avatar_url":"https://avatars.githubusercontent.com/u/146881603?v=4","primary_language":"Python","stars":358,"forks":38,"topics":["alignment","llm","llm-finetuning"],"archived":false,"github_pushed_at":"2024-02-23T21:19:44+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/llm-tuning-safety-llms-finetuning-safety","markdown_url":"https://www.graphcanon.com/tools/llm-tuning-safety-llms-finetuning-safety.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/llm-tuning-safety-llms-finetuning-safety","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=llm-tuning-safety-llms-finetuning-safety","shared_categories":["evaluation-observability"]},{"slug":"libr-ai-do-not-answer","name":"do-not-answer","tagline":"A Dataset for Evaluating Safeguards in LLMs","github_url":"https://github.com/Libr-AI/do-not-answer","owner":"Libr-AI","repo":"do-not-answer","owner_avatar_url":"https://avatars.githubusercontent.com/u/133515165?v=4","primary_language":"Jupyter Notebook","stars":343,"forks":29,"topics":[],"archived":false,"github_pushed_at":"2024-06-07T14:55:51+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/libr-ai-do-not-answer","markdown_url":"https://www.graphcanon.com/tools/libr-ai-do-not-answer.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/libr-ai-do-not-answer","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=libr-ai-do-not-answer","shared_categories":["evaluation-observability"]},{"slug":"cvs-health-langfair","name":"langfair","tagline":"LangFair: Use-Case Level LLM Bias and Fairness Assessments","github_url":"https://github.com/cvs-health/langfair","owner":"cvs-health","repo":"langfair","owner_avatar_url":"https://avatars.githubusercontent.com/u/89528242?v=4","primary_language":"Python","stars":261,"forks":46,"topics":["ai","ai-safety","artificial-intelligence","bias","bias-detection","ethical-ai","fairness","fairness-ai","fairness-ml","fairness-testing","large-language-models","llm","llm-evaluation","llm-evaluation-framework","llm-evaluation-metrics","python","responsible-ai"],"archived":false,"github_pushed_at":"2026-08-26T20:46:22+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/cvs-health-langfair","markdown_url":"https://www.graphcanon.com/tools/cvs-health-langfair.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/cvs-health-langfair","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=cvs-health-langfair","shared_categories":["evaluation-observability"]},{"slug":"microsoft-bipia","name":"BIPIA","tagline":"Benchmark for evaluating LLM robustness to indirect prompt injection attacks.","github_url":"https://github.com/microsoft/BIPIA","owner":"microsoft","repo":"BIPIA","owner_avatar_url":"https://avatars.githubusercontent.com/u/6154722?v=4","primary_language":"Python","stars":156,"forks":20,"topics":["llm-security"],"archived":false,"github_pushed_at":"2024-04-15T02:08:17+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/microsoft-bipia","markdown_url":"https://www.graphcanon.com/tools/microsoft-bipia.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/microsoft-bipia","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=microsoft-bipia","shared_categories":["evaluation-observability"]}]}}