{"data":{"node":{"slug":"b4rtaz-distributed-llama","name":"distributed-llama","tagline":"Distributed LLM inference using home devices cluster","github_url":"https://github.com/b4rtaz/distributed-llama","owner":"b4rtaz","repo":"distributed-llama","owner_avatar_url":"https://avatars.githubusercontent.com/u/12797776?v=4","primary_language":"C++","stars":3044,"forks":246,"topics":["distributed-computing","distributed-llm","llama2","llama3","llm","llm-inference","llms","neural-network","open-llm"],"archived":false,"github_pushed_at":"2026-07-05T16:47:20+00:00","maintenance_label":"Steady","stars_delta_30d":32,"url":"https://www.graphcanon.com/tools/b4rtaz-distributed-llama","markdown_url":"https://www.graphcanon.com/tools/b4rtaz-distributed-llama.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/b4rtaz-distributed-llama","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=b4rtaz-distributed-llama"},"categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"}],"tags":[{"slug":"distributed-computing","name":"distributed-computing"},{"slug":"llm-inference","name":"llm-inference"},{"slug":"neural-network","name":"neural-network"}],"edges":[],"neighbours":[{"slug":"alexsjones-llmfit","name":"llmfit","tagline":"Hundreds of models & providers. One command to find what runs on your hardware.","github_url":"https://github.com/AlexsJones/llmfit","owner":"AlexsJones","repo":"llmfit","owner_avatar_url":"https://avatars.githubusercontent.com/u/1235925?v=4","primary_language":"Rust","stars":31867,"forks":1978,"topics":["gguf","llm","localai","mlx","skill","unsloth"],"archived":false,"github_pushed_at":"2026-08-14T07:36:41+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/alexsjones-llmfit","markdown_url":"https://www.graphcanon.com/tools/alexsjones-llmfit.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/alexsjones-llmfit","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=alexsjones-llmfit","shared_categories":[]},{"slug":"steven2358-awesome-generative-ai","name":"awesome-generative-ai","tagline":"A curated list of modern Generative Artificial Intelligence projects and services","github_url":"https://github.com/steven2358/awesome-generative-ai","owner":"steven2358","repo":"awesome-generative-ai","owner_avatar_url":"https://avatars.githubusercontent.com/u/164072?v=4","primary_language":null,"stars":12501,"forks":1990,"topics":["ai","artificial-intelligence","awesome","awesome-list","generative-ai","generative-art","large-language-models","llm"],"archived":false,"github_pushed_at":"2026-08-03T10:58:05+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/steven2358-awesome-generative-ai","markdown_url":"https://www.graphcanon.com/tools/steven2358-awesome-generative-ai.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/steven2358-awesome-generative-ai","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=steven2358-awesome-generative-ai","shared_categories":["inference-serving"]},{"slug":"andyyyy64-whichllm","name":"whichllm","tagline":"Command-line tool to find and benchmark local LLM performance","github_url":"https://github.com/Andyyyy64/whichllm","owner":"Andyyyy64","repo":"whichllm","owner_avatar_url":"https://avatars.githubusercontent.com/u/105579829?v=4","primary_language":"Python","stars":6225,"forks":330,"topics":["ai","apple-silicon","benchmarks","cli","command-line-tool","gguf","gpu","huggingface","inference","llm","local-llm","ollama","python","vram"],"archived":false,"github_pushed_at":"2026-08-05T07:15:32+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/andyyyy64-whichllm","markdown_url":"https://www.graphcanon.com/tools/andyyyy64-whichllm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/andyyyy64-whichllm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=andyyyy64-whichllm","shared_categories":["inference-serving"]},{"slug":"osmantic-ods","name":"ODS","tagline":"Transform personal computers into AI servers.","github_url":"https://github.com/Osmantic/ODS","owner":"Osmantic","repo":"ODS","owner_avatar_url":"https://avatars.githubusercontent.com/u/262014141?v=4","primary_language":"Python","stars":3799,"forks":551,"topics":["ai-agents","amd","comfyui","docker","llama-cpp","llm","local-ai","n8n","nvidia","open-webui","rag","self-hosted","speech-to-text","strix-halo","text-to-speech","workflow-automation"],"archived":false,"github_pushed_at":"2026-07-29T11:57:22+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/osmantic-ods","markdown_url":"https://www.graphcanon.com/tools/osmantic-ods.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/osmantic-ods","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=osmantic-ods","shared_categories":["inference-serving"]},{"slug":"superlinked-sie","name":"sie","tagline":"Open-source inference server and production cluster for all the models your agent needs.","github_url":"https://github.com/superlinked/sie","owner":"superlinked","repo":"sie","owner_avatar_url":"https://avatars.githubusercontent.com/u/94243920?v=4","primary_language":"Python","stars":2804,"forks":272,"topics":["bge","colbert","data-pipeline","deep-learning","embeddings","inference","inference-server","information-retrieval","llm","ml","mlops","natural-language-processing","nlp","python","reranking","retrieval","retrieval-augmented-generation","semantic-search","splade","vector-search"],"archived":false,"github_pushed_at":"2026-08-21T20:28:04+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/superlinked-sie","markdown_url":"https://www.graphcanon.com/tools/superlinked-sie.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/superlinked-sie","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=superlinked-sie","shared_categories":["inference-serving"]},{"slug":"rafska-awesome-local-llm","name":"awesome-local-llm","tagline":"Resources for running LLMs locally","github_url":"https://github.com/rafska/awesome-local-llm","owner":"rafska","repo":"awesome-local-llm","owner_avatar_url":"https://avatars.githubusercontent.com/u/17859377?v=4","primary_language":null,"stars":2518,"forks":316,"topics":["ai","awesome","awesome-list","llm","local","local-ai","local-llm","resources","self-hosted","selfhosted"],"archived":false,"github_pushed_at":"2026-08-04T23:08:15+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/rafska-awesome-local-llm","markdown_url":"https://www.graphcanon.com/tools/rafska-awesome-local-llm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/rafska-awesome-local-llm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=rafska-awesome-local-llm","shared_categories":["inference-serving"]},{"slug":"theopenco-llmgateway","name":"llmgateway","tagline":"Route and manage LLM requests via unified API interface","github_url":"https://github.com/theopenco/llmgateway","owner":"theopenco","repo":"llmgateway","owner_avatar_url":"https://avatars.githubusercontent.com/u/211671860?v=4","primary_language":"TypeScript","stars":1518,"forks":170,"topics":["ai","ai-gateway","analytics","anthropic","api-key-management","claude","codex","enterprise","guardrails","inference","llm","llm-gateway","llm-proxy","llms","observability","openai","opencode","rate-limiting","typescript"],"archived":false,"github_pushed_at":"2026-08-09T00:38:36+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/theopenco-llmgateway","markdown_url":"https://www.graphcanon.com/tools/theopenco-llmgateway.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/theopenco-llmgateway","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=theopenco-llmgateway","shared_categories":["inference-serving"]},{"slug":"waybarrios-vllm-mlx","name":"vllm-mlx","tagline":"Server for LLMs and vision-language models compatible with Apple Silicon","github_url":"https://github.com/waybarrios/vllm-mlx","owner":"waybarrios","repo":"vllm-mlx","owner_avatar_url":"https://avatars.githubusercontent.com/u/6794828?v=4","primary_language":"Python","stars":1472,"forks":205,"topics":["anthropic","apple-silicon","audio-processing","claude-code","computer-vision","image-understanding","inference","llm","machine-learning","macos","mllm","mlx","multimodal-ai","speech-to-text","stt","text-to-speech","tts","video-understanding","vision-language-model","vllm"],"archived":false,"github_pushed_at":"2026-06-28T20:18:31+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/waybarrios-vllm-mlx","markdown_url":"https://www.graphcanon.com/tools/waybarrios-vllm-mlx.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/waybarrios-vllm-mlx","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=waybarrios-vllm-mlx","shared_categories":["inference-serving"]},{"slug":"atomicbot-ai-atomic-chat","name":"Atomic-Chat","tagline":"Local AI app and inference engine for agents","github_url":"https://github.com/AtomicBot-ai/Atomic-Chat","owner":"AtomicBot-ai","repo":"Atomic-Chat","owner_avatar_url":"https://avatars.githubusercontent.com/u/259533419?v=4","primary_language":"TypeScript","stars":1363,"forks":157,"topics":["ai-chat","ai-tools","apple-silicon","chatgpt","deepseek","desktop-app","gemma","gguf","gpt-oss","llamacpp","llm","llm-inference","local-ai","local-first","local-llm","mcp","mlx","open-source","qwen","self-hosted"],"archived":false,"github_pushed_at":"2026-08-24T14:20:33+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/atomicbot-ai-atomic-chat","markdown_url":"https://www.graphcanon.com/tools/atomicbot-ai-atomic-chat.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/atomicbot-ai-atomic-chat","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=atomicbot-ai-atomic-chat","shared_categories":["inference-serving"]},{"slug":"kubeai-project-kubeai","name":"kubeai","tagline":"AI Inference Operator for Kubernetes","github_url":"https://github.com/kubeai-project/kubeai","owner":"kubeai-project","repo":"kubeai","owner_avatar_url":"https://avatars.githubusercontent.com/u/232319222?v=4","primary_language":"Go","stars":1237,"forks":131,"topics":["ai","autoscaler","faster-whisper","inference-operator","k8s","kubernetes","llm","ollama","ollama-operator","openai-api","vllm","vllm-operator","whisper"],"archived":false,"github_pushed_at":"2026-07-31T01:04:47+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/kubeai-project-kubeai","markdown_url":"https://www.graphcanon.com/tools/kubeai-project-kubeai.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/kubeai-project-kubeai","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=kubeai-project-kubeai","shared_categories":["inference-serving"]},{"slug":"openinfer-project-openinfer","name":"openinfer","tagline":"Pure Rust CUDA LLM inference engine serving multiple models including Qwen3 and Kimi-K2","github_url":"https://github.com/openinfer-project/openinfer","owner":"openinfer-project","repo":"openinfer","owner_avatar_url":"https://avatars.githubusercontent.com/u/292134277?v=4","primary_language":"Rust","stars":657,"forks":103,"topics":["cuda","cuda-kernels","deepseek","gpu","inference","inference-engine","kimi","kimi-k2","kv-cache","llm","llm-inference","llm-serving","model-serving","moe","openai-api","paged-attention","qwen","qwen3","rust","vllm"],"archived":false,"github_pushed_at":"2026-08-25T05:44:30+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/openinfer-project-openinfer","markdown_url":"https://www.graphcanon.com/tools/openinfer-project-openinfer.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/openinfer-project-openinfer","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=openinfer-project-openinfer","shared_categories":["inference-serving"]},{"slug":"timmyy123-llm-hub","name":"LLM-Hub","tagline":"Local AI Assistant on your phone","github_url":"https://github.com/timmyy123/LLM-Hub","owner":"timmyy123","repo":"LLM-Hub","owner_avatar_url":"https://avatars.githubusercontent.com/u/164847362?v=4","primary_language":"C++","stars":561,"forks":116,"topics":["ai","gemma","gemma4","gemma4-agent-skills","gptoss","granite","imagegeneration","lfm25","llama","llm","llm-inference","mistral","music","musicgeneration","phi4","rag","stable-diffusion","videogeneration","whisper"],"archived":false,"github_pushed_at":"2026-08-24T07:20:33+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/timmyy123-llm-hub","markdown_url":"https://www.graphcanon.com/tools/timmyy123-llm-hub.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/timmyy123-llm-hub","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=timmyy123-llm-hub","shared_categories":["inference-serving"]}]}}