{"data":{"node":{"slug":"kubeai-project-kubeai","name":"kubeai","tagline":"AI Inference Operator for Kubernetes","github_url":"https://github.com/kubeai-project/kubeai","owner":"kubeai-project","repo":"kubeai","owner_avatar_url":"https://avatars.githubusercontent.com/u/232319222?v=4","primary_language":"Go","stars":1237,"forks":131,"topics":["ai","autoscaler","faster-whisper","inference-operator","k8s","kubernetes","llm","ollama","ollama-operator","openai-api","vllm","vllm-operator","whisper"],"archived":false,"github_pushed_at":"2026-07-31T01:04:47+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/kubeai-project-kubeai","markdown_url":"https://www.graphcanon.com/tools/kubeai-project-kubeai.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/kubeai-project-kubeai","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=kubeai-project-kubeai"},"categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"},{"slug":"llm-frameworks","name":"LLM Frameworks","url":"https://www.graphcanon.com/categories/llm-frameworks","markdown_url":"https://www.graphcanon.com/categories/llm-frameworks.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/llm-frameworks"},{"slug":"speech-audio","name":"Speech & Audio","url":"https://www.graphcanon.com/categories/speech-audio","markdown_url":"https://www.graphcanon.com/categories/speech-audio.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/speech-audio"}],"tags":[{"slug":"ai","name":"ai"},{"slug":"autoscaler","name":"autoscaler"},{"slug":"faster-whisper","name":"faster-whisper"},{"slug":"inference-operator","name":"inference-operator"},{"slug":"k8s","name":"k8s"},{"slug":"kubernetes","name":"kubernetes"},{"slug":"llm","name":"llm"},{"slug":"ollama","name":"ollama"}],"edges":[],"neighbours":[{"slug":"alexsjones-llmfit","name":"llmfit","tagline":"Hundreds of models & providers. One command to find what runs on your hardware.","github_url":"https://github.com/AlexsJones/llmfit","owner":"AlexsJones","repo":"llmfit","owner_avatar_url":"https://avatars.githubusercontent.com/u/1235925?v=4","primary_language":"Rust","stars":31867,"forks":1978,"topics":["gguf","llm","localai","mlx","skill","unsloth"],"archived":false,"github_pushed_at":"2026-08-14T07:36:41+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/alexsjones-llmfit","markdown_url":"https://www.graphcanon.com/tools/alexsjones-llmfit.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/alexsjones-llmfit","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=alexsjones-llmfit","shared_categories":["llm-frameworks"]},{"slug":"steven2358-awesome-generative-ai","name":"awesome-generative-ai","tagline":"A curated list of modern Generative Artificial Intelligence projects and services","github_url":"https://github.com/steven2358/awesome-generative-ai","owner":"steven2358","repo":"awesome-generative-ai","owner_avatar_url":"https://avatars.githubusercontent.com/u/164072?v=4","primary_language":null,"stars":12501,"forks":1990,"topics":["ai","artificial-intelligence","awesome","awesome-list","generative-ai","generative-art","large-language-models","llm"],"archived":false,"github_pushed_at":"2026-08-03T10:58:05+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/steven2358-awesome-generative-ai","markdown_url":"https://www.graphcanon.com/tools/steven2358-awesome-generative-ai.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/steven2358-awesome-generative-ai","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=steven2358-awesome-generative-ai","shared_categories":["llm-frameworks","inference-serving"]},{"slug":"oumi-ai-oumi","name":"oumi","tagline":"Easily fine-tune, evaluate and deploy open source LLMs/VLMs","github_url":"https://github.com/oumi-ai/oumi","owner":"oumi-ai","repo":"oumi","owner_avatar_url":"https://avatars.githubusercontent.com/u/167452922?v=4","primary_language":"Python","stars":9376,"forks":784,"topics":["dpo","evaluation","fine-tuning","gpt-oss","gpt-oss-120b","gpt-oss-20b","inference","llama","llms","open-weight","open-weight-models","open-weights","sft","slms","vlms"],"archived":false,"github_pushed_at":"2026-08-21T23:11:35+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/oumi-ai-oumi","markdown_url":"https://www.graphcanon.com/tools/oumi-ai-oumi.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/oumi-ai-oumi","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=oumi-ai-oumi","shared_categories":["inference-serving"]},{"slug":"andyyyy64-whichllm","name":"whichllm","tagline":"Command-line tool to find and benchmark local LLM performance","github_url":"https://github.com/Andyyyy64/whichllm","owner":"Andyyyy64","repo":"whichllm","owner_avatar_url":"https://avatars.githubusercontent.com/u/105579829?v=4","primary_language":"Python","stars":6225,"forks":330,"topics":["ai","apple-silicon","benchmarks","cli","command-line-tool","gguf","gpu","huggingface","inference","llm","local-llm","ollama","python","vram"],"archived":false,"github_pushed_at":"2026-08-05T07:15:32+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/andyyyy64-whichllm","markdown_url":"https://www.graphcanon.com/tools/andyyyy64-whichllm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/andyyyy64-whichllm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=andyyyy64-whichllm","shared_categories":["inference-serving"]},{"slug":"huaizhengzhang-ai-infra-from-zero-to-hero","name":"AI-Infra-from-Zero-to-Hero","tagline":"Awesome System for Machine Learning and LLM Infra","github_url":"https://github.com/HuaizhengZhang/AI-Infra-from-Zero-to-Hero","owner":"HuaizhengZhang","repo":"AI-Infra-from-Zero-to-Hero","owner_avatar_url":"https://avatars.githubusercontent.com/u/5894780?v=4","primary_language":null,"stars":4285,"forks":409,"topics":["ai-infra","genai","large-language-models","llmsys","mlsys","model-serving","model-training"],"archived":false,"github_pushed_at":"2025-07-25T02:24:35+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/huaizhengzhang-ai-infra-from-zero-to-hero","markdown_url":"https://www.graphcanon.com/tools/huaizhengzhang-ai-infra-from-zero-to-hero.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/huaizhengzhang-ai-infra-from-zero-to-hero","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=huaizhengzhang-ai-infra-from-zero-to-hero","shared_categories":["llm-frameworks","inference-serving"]},{"slug":"nyldn-claude-octopus","name":"claude-octopus","tagline":"Surface AI blindspots before you ship","github_url":"https://github.com/nyldn/claude-octopus","owner":"nyldn","repo":"claude-octopus","owner_avatar_url":"https://avatars.githubusercontent.com/u/4805949?v=4","primary_language":"Shell","stars":3962,"forks":374,"topics":["ai-agents","ai-orchestration","claude-code","claude-code-plugin","codex","copilot","developer-tools","double-diamond","gemini","multi-ai","multi-llm","ollama"],"archived":false,"github_pushed_at":"2026-08-13T23:28:23+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/nyldn-claude-octopus","markdown_url":"https://www.graphcanon.com/tools/nyldn-claude-octopus.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/nyldn-claude-octopus","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=nyldn-claude-octopus","shared_categories":[]},{"slug":"osmantic-ods","name":"ODS","tagline":"Transform personal computers into AI servers.","github_url":"https://github.com/Osmantic/ODS","owner":"Osmantic","repo":"ODS","owner_avatar_url":"https://avatars.githubusercontent.com/u/262014141?v=4","primary_language":"Python","stars":3799,"forks":551,"topics":["ai-agents","amd","comfyui","docker","llama-cpp","llm","local-ai","n8n","nvidia","open-webui","rag","self-hosted","speech-to-text","strix-halo","text-to-speech","workflow-automation"],"archived":false,"github_pushed_at":"2026-07-29T11:57:22+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/osmantic-ods","markdown_url":"https://www.graphcanon.com/tools/osmantic-ods.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/osmantic-ods","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=osmantic-ods","shared_categories":["speech-audio","inference-serving"]},{"slug":"b4rtaz-distributed-llama","name":"distributed-llama","tagline":"Distributed LLM inference using home devices cluster","github_url":"https://github.com/b4rtaz/distributed-llama","owner":"b4rtaz","repo":"distributed-llama","owner_avatar_url":"https://avatars.githubusercontent.com/u/12797776?v=4","primary_language":"C++","stars":3044,"forks":246,"topics":["distributed-computing","distributed-llm","llama2","llama3","llm","llm-inference","llms","neural-network","open-llm"],"archived":false,"github_pushed_at":"2026-07-05T16:47:20+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/b4rtaz-distributed-llama","markdown_url":"https://www.graphcanon.com/tools/b4rtaz-distributed-llama.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/b4rtaz-distributed-llama","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=b4rtaz-distributed-llama","shared_categories":["inference-serving"]},{"slug":"superlinked-sie","name":"sie","tagline":"Open-source inference server and production cluster for all the models your agent needs.","github_url":"https://github.com/superlinked/sie","owner":"superlinked","repo":"sie","owner_avatar_url":"https://avatars.githubusercontent.com/u/94243920?v=4","primary_language":"Python","stars":2804,"forks":272,"topics":["bge","colbert","data-pipeline","deep-learning","embeddings","inference","inference-server","information-retrieval","llm","ml","mlops","natural-language-processing","nlp","python","reranking","retrieval","retrieval-augmented-generation","semantic-search","splade","vector-search"],"archived":false,"github_pushed_at":"2026-08-21T20:28:04+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/superlinked-sie","markdown_url":"https://www.graphcanon.com/tools/superlinked-sie.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/superlinked-sie","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=superlinked-sie","shared_categories":["inference-serving"]},{"slug":"waybarrios-vllm-mlx","name":"vllm-mlx","tagline":"Server for LLMs and vision-language models compatible with Apple Silicon","github_url":"https://github.com/waybarrios/vllm-mlx","owner":"waybarrios","repo":"vllm-mlx","owner_avatar_url":"https://avatars.githubusercontent.com/u/6794828?v=4","primary_language":"Python","stars":1472,"forks":205,"topics":["anthropic","apple-silicon","audio-processing","claude-code","computer-vision","image-understanding","inference","llm","machine-learning","macos","mllm","mlx","multimodal-ai","speech-to-text","stt","text-to-speech","tts","video-understanding","vision-language-model","vllm"],"archived":false,"github_pushed_at":"2026-06-28T20:18:31+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/waybarrios-vllm-mlx","markdown_url":"https://www.graphcanon.com/tools/waybarrios-vllm-mlx.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/waybarrios-vllm-mlx","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=waybarrios-vllm-mlx","shared_categories":["inference-serving"]},{"slug":"ebhy-budgetml","name":"budgetml","tagline":"Deploys ML inference service economically","github_url":"https://github.com/ebhy/budgetml","owner":"ebhy","repo":"budgetml","owner_avatar_url":"https://avatars.githubusercontent.com/u/76654256?v=4","primary_language":"Python","stars":1343,"forks":65,"topics":["api","data-science","deployment","fastapi","inference","machine-learning","mlops"],"archived":false,"github_pushed_at":"2024-02-12T17:29:24+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/ebhy-budgetml","markdown_url":"https://www.graphcanon.com/tools/ebhy-budgetml.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ebhy-budgetml","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ebhy-budgetml","shared_categories":["inference-serving"]},{"slug":"daoyuanli2816-can-i-finetune-this","name":"can-i-finetune-this","tagline":"Estimate if a Hugging Face model can fine-tune locally on GPU","github_url":"https://github.com/DaoyuanLi2816/can-i-finetune-this","owner":"DaoyuanLi2816","repo":"can-i-finetune-this","owner_avatar_url":"https://avatars.githubusercontent.com/u/94409450?v=4","primary_language":"Python","stars":792,"forks":107,"topics":["bitsandbytes","fine-tuning","gpu","hugging-face","llm","lora","memory-estimation","peft","pytorch","qlora","transformers","vram"],"archived":false,"github_pushed_at":"2026-07-23T05:11:25+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/daoyuanli2816-can-i-finetune-this","markdown_url":"https://www.graphcanon.com/tools/daoyuanli2816-can-i-finetune-this.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/daoyuanli2816-can-i-finetune-this","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=daoyuanli2816-can-i-finetune-this","shared_categories":["llm-frameworks"]}]}}