{"data":{"node":{"slug":"xamey-deploy-llms-with-ansible","name":"deploy-llms-with-ansible","tagline":"Easily deploy LLMs using Ansible","github_url":"https://github.com/xamey/deploy-llms-with-ansible","owner":"xamey","repo":"deploy-llms-with-ansible","owner_avatar_url":"https://avatars.githubusercontent.com/u/34269296?v=4","primary_language":null,"stars":3,"forks":0,"topics":[],"archived":false,"github_pushed_at":"2025-05-01T21:58:25+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/xamey-deploy-llms-with-ansible","markdown_url":"https://www.graphcanon.com/tools/xamey-deploy-llms-with-ansible.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/xamey-deploy-llms-with-ansible","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=xamey-deploy-llms-with-ansible"},"categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"}],"tags":[{"slug":"ansible","name":"ansible"},{"slug":"deployment","name":"deployment"},{"slug":"docker","name":"docker"},{"slug":"llama-cpp","name":"llama-cpp"},{"slug":"ollama","name":"ollama"},{"slug":"vm","name":"vm"},{"slug":"whitelisting","name":"whitelisting"}],"edges":[],"neighbours":[{"slug":"vllm-project-vllm","name":"vllm","tagline":"A high-throughput and memory-efficient inference and serving engine for LLMs","github_url":"https://github.com/vllm-project/vllm","owner":"vllm-project","repo":"vllm","owner_avatar_url":"https://avatars.githubusercontent.com/u/136984999?v=4","primary_language":"Python","stars":87847,"forks":20135,"topics":["amd","blackwell","cuda","deepseek","deepseek-v3","gpt","gpt-oss","inference","kimi","llama","llm","llm-serving","model-serving","moe","openai","pytorch","qwen","qwen3","tpu","transformer"],"archived":false,"github_pushed_at":"2026-08-01T11:55:36+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/vllm-project-vllm","markdown_url":"https://www.graphcanon.com/tools/vllm-project-vllm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/vllm-project-vllm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=vllm-project-vllm","shared_categories":["inference-serving"]},{"slug":"lightning-ai-litgpt","name":"litgpt","tagline":"High-performance LLMs with recipes for pretraining, finetuning and deployment","github_url":"https://github.com/Lightning-AI/litgpt","owner":"Lightning-AI","repo":"litgpt","owner_avatar_url":"https://avatars.githubusercontent.com/u/58386951?v=4","primary_language":"Python","stars":13605,"forks":1483,"topics":["ai","artificial-intelligence","deep-learning","large-language-models","llm","llm-inference","llms"],"archived":false,"github_pushed_at":"2026-07-20T10:24:12+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/lightning-ai-litgpt","markdown_url":"https://www.graphcanon.com/tools/lightning-ai-litgpt.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/lightning-ai-litgpt","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=lightning-ai-litgpt","shared_categories":["inference-serving"]},{"slug":"rafska-awesome-local-llm","name":"awesome-local-llm","tagline":"Resources for running LLMs locally","github_url":"https://github.com/rafska/awesome-local-llm","owner":"rafska","repo":"awesome-local-llm","owner_avatar_url":"https://avatars.githubusercontent.com/u/17859377?v=4","primary_language":null,"stars":2518,"forks":316,"topics":["ai","awesome","awesome-list","llm","local","local-ai","local-llm","resources","self-hosted","selfhosted"],"archived":false,"github_pushed_at":"2026-08-04T23:08:15+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/rafska-awesome-local-llm","markdown_url":"https://www.graphcanon.com/tools/rafska-awesome-local-llm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/rafska-awesome-local-llm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=rafska-awesome-local-llm","shared_categories":["inference-serving"]},{"slug":"jhubi1-ollama-app","name":"ollama-app","tagline":"A modern and easy-to-use client for Ollama","github_url":"https://github.com/JHubi1/ollama-app","owner":"JHubi1","repo":"ollama-app","owner_avatar_url":"https://avatars.githubusercontent.com/u/61345690?v=4","primary_language":"Dart","stars":1797,"forks":216,"topics":["ai","android","app","linux","llama","localai","ollama","ollama-client","windows"],"archived":false,"github_pushed_at":"2026-08-07T21:26:25+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/jhubi1-ollama-app","markdown_url":"https://www.graphcanon.com/tools/jhubi1-ollama-app.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/jhubi1-ollama-app","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=jhubi1-ollama-app","shared_categories":[]},{"slug":"waybarrios-vllm-mlx","name":"vllm-mlx","tagline":"Server for LLMs and vision-language models compatible with Apple Silicon","github_url":"https://github.com/waybarrios/vllm-mlx","owner":"waybarrios","repo":"vllm-mlx","owner_avatar_url":"https://avatars.githubusercontent.com/u/6794828?v=4","primary_language":"Python","stars":1472,"forks":205,"topics":["anthropic","apple-silicon","audio-processing","claude-code","computer-vision","image-understanding","inference","llm","machine-learning","macos","mllm","mlx","multimodal-ai","speech-to-text","stt","text-to-speech","tts","video-understanding","vision-language-model","vllm"],"archived":false,"github_pushed_at":"2026-06-28T20:18:31+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/waybarrios-vllm-mlx","markdown_url":"https://www.graphcanon.com/tools/waybarrios-vllm-mlx.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/waybarrios-vllm-mlx","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=waybarrios-vllm-mlx","shared_categories":["inference-serving"]},{"slug":"benman1-generative-ai-with-langchain","name":"generative_ai_with_langchain","tagline":"Build production-ready LLM applications and advanced agents using Python, LangChain, and LangGraph","github_url":"https://github.com/benman1/generative_ai_with_langchain","owner":"benman1","repo":"generative_ai_with_langchain","owner_avatar_url":"https://avatars.githubusercontent.com/u/10786684?v=4","primary_language":"Jupyter Notebook","stars":1400,"forks":582,"topics":["agent","chatgpt","claude","claude-3-5-sonnet","deepseek","deepseek-r1","gpt","gpt-4o","huggingface","langchain","langgraph","llamacpp","llms","ollama","openai"],"archived":false,"github_pushed_at":"2026-08-05T12:50:30+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/benman1-generative-ai-with-langchain","markdown_url":"https://www.graphcanon.com/tools/benman1-generative-ai-with-langchain.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/benman1-generative-ai-with-langchain","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=benman1-generative-ai-with-langchain","shared_categories":[]},{"slug":"oxbshw-llm-agents-ecosystem-handbook","name":"LLM-Agents-Ecosystem-Handbook","tagline":"One-stop handbook for building, deploying, and understanding LLM agents","github_url":"https://github.com/oxbshw/LLM-Agents-Ecosystem-Handbook","owner":"oxbshw","repo":"LLM-Agents-Ecosystem-Handbook","owner_avatar_url":"https://avatars.githubusercontent.com/u/212214682?v=4","primary_language":"Python","stars":539,"forks":85,"topics":["ai","ai-agent","ai-agents","fine-tuning","finetuning-llms","freamework","llm","llmops","local-development","mcp-server","memory","rag","rag-chatbot","voice-agent"],"archived":false,"github_pushed_at":"2026-06-30T12:22:57+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/oxbshw-llm-agents-ecosystem-handbook","markdown_url":"https://www.graphcanon.com/tools/oxbshw-llm-agents-ecosystem-handbook.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/oxbshw-llm-agents-ecosystem-handbook","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=oxbshw-llm-agents-ecosystem-handbook","shared_categories":[]},{"slug":"kaito-project-aikit","name":"aikit","tagline":"Fine-tune, build, and deploy open-source LLMs easily!","github_url":"https://github.com/kaito-project/aikit","owner":"kaito-project","repo":"aikit","owner_avatar_url":"https://avatars.githubusercontent.com/u/186863079?v=4","primary_language":"Go","stars":537,"forks":57,"topics":["ai","buildkit","chatgpt","docker","fine-tuning","finetuning","gemma","gpt","inference","kubernetes","large-language-models","llama","llm","localllama","mistral","mixtral","nvidia","open-llm","open-source-llm","openai"],"archived":false,"github_pushed_at":"2026-08-24T02:38:26+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/kaito-project-aikit","markdown_url":"https://www.graphcanon.com/tools/kaito-project-aikit.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/kaito-project-aikit","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=kaito-project-aikit","shared_categories":["inference-serving"]},{"slug":"chen-zexi-vllm-cli","name":"vllm-cli","tagline":"Command-line interface for serving LLM using vLLM","github_url":"https://github.com/Chen-zexi/vllm-cli","owner":"Chen-zexi","repo":"vllm-cli","owner_avatar_url":"https://avatars.githubusercontent.com/u/128259419?v=4","primary_language":"Python","stars":506,"forks":29,"topics":["llm","llm-inference","llm-tools","vllm"],"archived":false,"github_pushed_at":"2026-01-25T19:37:43+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/chen-zexi-vllm-cli","markdown_url":"https://www.graphcanon.com/tools/chen-zexi-vllm-cli.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/chen-zexi-vllm-cli","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=chen-zexi-vllm-cli","shared_categories":["inference-serving"]},{"slug":"anarchy-ai-llm-vm","name":"LLM-VM","tagline":"irresponsible innovation","github_url":"https://github.com/anarchy-ai/LLM-VM","owner":"anarchy-ai","repo":"LLM-VM","owner_avatar_url":"https://avatars.githubusercontent.com/u/134051110?v=4","primary_language":"Python","stars":490,"forks":139,"topics":["artificial-intelligence","deep-learning","distillation","distillation-model","llm","llm-agent","llm-inference","llm-local","llm-training","machine-learning"],"archived":false,"github_pushed_at":"2024-05-14T07:38:07+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/anarchy-ai-llm-vm","markdown_url":"https://www.graphcanon.com/tools/anarchy-ai-llm-vm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/anarchy-ai-llm-vm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=anarchy-ai-llm-vm","shared_categories":["inference-serving"]},{"slug":"tensorchord-openmodelz","name":"openmodelz","tagline":"Automate and scale inference of large language models on Kubernetes.","github_url":"https://github.com/tensorchord/openmodelz","owner":"tensorchord","repo":"openmodelz","owner_avatar_url":"https://avatars.githubusercontent.com/u/100543303?v=4","primary_language":"Go","stars":282,"forks":26,"topics":["cluster-manager","hacktoberfest","inference","llm","llmops","mlops"],"archived":false,"github_pushed_at":"2023-11-03T06:33:25+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/tensorchord-openmodelz","markdown_url":"https://www.graphcanon.com/tools/tensorchord-openmodelz.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/tensorchord-openmodelz","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=tensorchord-openmodelz","shared_categories":["inference-serving"]},{"slug":"mohitsoni48-turbollm","name":"TurboLLM","tagline":"Run any local LLM engine auto-tuned to your GPU with polished web UI and OpenAI/Anthropic-compatible API","github_url":"https://github.com/mohitsoni48/TurboLLM","owner":"mohitsoni48","repo":"TurboLLM","owner_avatar_url":"https://avatars.githubusercontent.com/u/63787789?v=4","primary_language":"TypeScript","stars":225,"forks":36,"topics":["ai","anthropic-api","claude-code","gguf","gpu","inference","llama-cpp","llama-server","llm","local-llm","offline","openai-api","self-hosted"],"archived":false,"github_pushed_at":"2026-08-11T13:28:32+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/mohitsoni48-turbollm","markdown_url":"https://www.graphcanon.com/tools/mohitsoni48-turbollm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/mohitsoni48-turbollm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=mohitsoni48-turbollm","shared_categories":["inference-serving"]}]}}