{"data":{"node":{"slug":"peva3-smarterrouter","name":"SmarterRouter","tagline":"An intelligent LLM gateway and VRAM-aware router for Ollama, llama.cpp, and OpenAI.","github_url":"https://github.com/peva3/SmarterRouter","owner":"peva3","repo":"SmarterRouter","owner_avatar_url":"https://avatars.githubusercontent.com/u/1492185?v=4","primary_language":"Python","stars":152,"forks":19,"topics":["ai-cache","ai-gateway","docker","fastapi","gpu-monitoring","llm","llm-proxy","llm-router","local-llm","model-serving","ollama","ollama-api","openai-proxy","self-hosted","self-hosted-ai","semantic-cache"],"archived":false,"github_pushed_at":"2026-05-10T02:47:50+00:00","maintenance_label":"Slowing","stars_delta_30d":6,"url":"https://www.graphcanon.com/tools/peva3-smarterrouter","markdown_url":"https://www.graphcanon.com/tools/peva3-smarterrouter.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/peva3-smarterrouter","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=peva3-smarterrouter"},"categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"},{"slug":"llm-frameworks","name":"LLM Frameworks","url":"https://www.graphcanon.com/categories/llm-frameworks","markdown_url":"https://www.graphcanon.com/categories/llm-frameworks.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/llm-frameworks"}],"tags":[{"slug":"ai-cache","name":"ai-cache"},{"slug":"ai-gateway","name":"ai-gateway"},{"slug":"docker-compose","name":"docker-compose"},{"slug":"fastapi","name":"fastapi"},{"slug":"gpu-monitoring","name":"gpu-monitoring"},{"slug":"llm-router","name":"llm-router"},{"slug":"model-serving","name":"model-serving"},{"slug":"ollama-api","name":"ollama-api"}],"edges":[],"neighbours":[{"slug":"open-webui-open-webui","name":"open-webui","tagline":"User-friendly AI Interface","github_url":"https://github.com/open-webui/open-webui","owner":"open-webui","repo":"open-webui","owner_avatar_url":"https://avatars.githubusercontent.com/u/158137808?v=4","primary_language":"Python","stars":152445,"forks":22306,"topics":["ai","llm","llm-ui","llm-webui","llms","mcp","ollama","ollama-webui","open-webui","openai","openapi","rag","self-hosted","ui","webui"],"archived":false,"github_pushed_at":"2026-09-18T00:11:06+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/open-webui-open-webui","markdown_url":"https://www.graphcanon.com/tools/open-webui-open-webui.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/open-webui-open-webui","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=open-webui-open-webui","shared_categories":["llm-frameworks","inference-serving"]},{"slug":"diegosouzapw-omniroute","name":"OmniRoute","tagline":"Free AI gateway with multi-provider support and token savings","github_url":"https://github.com/diegosouzapw/OmniRoute","owner":"diegosouzapw","repo":"OmniRoute","owner_avatar_url":"https://avatars.githubusercontent.com/u/8016841?v=4","primary_language":"TypeScript","stars":68316,"forks":9660,"topics":["a2a","ai-agents","ai-gateway","anthropic","claude","claude-code","cline","codex","copilot","cursor","deepseek","free-ai","gemini","kimi","llm-gateway","mcp","openai","openai-proxy","qwen","token-saver"],"archived":false,"github_pushed_at":"2026-09-19T08:30:04+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/diegosouzapw-omniroute","markdown_url":"https://www.graphcanon.com/tools/diegosouzapw-omniroute.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/diegosouzapw-omniroute","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=diegosouzapw-omniroute","shared_categories":["inference-serving"]},{"slug":"berriai-litellm","name":"litellm","tagline":"The fastest, lightest AI Gateway","github_url":"https://github.com/BerriAI/litellm","owner":"BerriAI","repo":"litellm","owner_avatar_url":"https://avatars.githubusercontent.com/u/121462774?v=4","primary_language":"Python","stars":59051,"forks":11547,"topics":["ai-gateway","anthropic","azure-openai","bedrock","gateway","langchain","litellm","llm","llm-gateway","llmops","mcp-gateway","openai","openai-proxy","rust","rust-ai","vertex-ai"],"archived":false,"github_pushed_at":"2026-09-18T08:06:17+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/berriai-litellm","markdown_url":"https://www.graphcanon.com/tools/berriai-litellm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/berriai-litellm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=berriai-litellm","shared_categories":["inference-serving"]},{"slug":"decolua-9router","name":"9router","tagline":"Unlimited FREE AI coding with auto-fallback and token-saving features","github_url":"https://github.com/decolua/9router","owner":"decolua","repo":"9router","owner_avatar_url":"https://avatars.githubusercontent.com/u/8282593?v=4","primary_language":"JavaScript","stars":29223,"forks":5390,"topics":["ai-agents","ai-gateway","anthropic","chatgpt","claude","claude-code","cline","codex","copilot","cursor","deepseek","free-ai","gemini","gemini-cli","llm","llm-gateway","openai","openai-proxy","qwen","token-saver"],"archived":false,"github_pushed_at":"2026-09-10T17:11:20+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/decolua-9router","markdown_url":"https://www.graphcanon.com/tools/decolua-9router.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/decolua-9router","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=decolua-9router","shared_categories":[]},{"slug":"portkey-ai-gateway","name":"gateway","tagline":"A high-performance AI Gateway connecting to over 1,600 LLMs with guardrails.","github_url":"https://github.com/Portkey-AI/gateway","owner":"Portkey-AI","repo":"gateway","owner_avatar_url":"https://avatars.githubusercontent.com/u/131141116?v=4","primary_language":"TypeScript","stars":12920,"forks":1293,"topics":["ai-gateway","gateway","generative-ai","hacktoberfest","langchain","llm","llm-gateway","llmops","llms","mcp","mcp-client","mcp-gateway","mcp-servers","model-router","openai"],"archived":false,"github_pushed_at":"2026-05-25T13:54:51+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/portkey-ai-gateway","markdown_url":"https://www.graphcanon.com/tools/portkey-ai-gateway.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/portkey-ai-gateway","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=portkey-ai-gateway","shared_categories":["llm-frameworks"]},{"slug":"andyyyy64-whichllm","name":"whichllm","tagline":"Command-line tool to find and benchmark local LLM performance","github_url":"https://github.com/Andyyyy64/whichllm","owner":"Andyyyy64","repo":"whichllm","owner_avatar_url":"https://avatars.githubusercontent.com/u/105579829?v=4","primary_language":"Python","stars":6666,"forks":368,"topics":["localllm"],"archived":false,"github_pushed_at":"2026-09-19T16:22:49+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/andyyyy64-whichllm","markdown_url":"https://www.graphcanon.com/tools/andyyyy64-whichllm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/andyyyy64-whichllm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=andyyyy64-whichllm","shared_categories":["inference-serving"]},{"slug":"vllm-project-semantic-router","name":"semantic-router","tagline":"System level intelligent runtime for Mixture-of-Models across edge, data center and cloud","github_url":"https://github.com/vllm-project/semantic-router","owner":"vllm-project","repo":"semantic-router","owner_avatar_url":"https://avatars.githubusercontent.com/u/136984999?v=4","primary_language":"Go","stars":5869,"forks":957,"topics":["ai-gateway","guardrails","inference","kubernetes","llm","llmrouter","mixture-of-models","pytorch","semantic-router","transformer","vllm"],"archived":false,"github_pushed_at":"2026-09-19T18:18:46+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/vllm-project-semantic-router","markdown_url":"https://www.graphcanon.com/tools/vllm-project-semantic-router.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/vllm-project-semantic-router","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=vllm-project-semantic-router","shared_categories":["llm-frameworks","inference-serving"]},{"slug":"aurelio-labs-semantic-router","name":"semantic-router","tagline":"Superfast AI decision making and intelligent processing of multi-modal data","github_url":"https://github.com/aurelio-labs/semantic-router","owner":"aurelio-labs","repo":"semantic-router","owner_avatar_url":"https://avatars.githubusercontent.com/u/69076224?v=4","primary_language":"Python","stars":3852,"forks":366,"topics":["ai","artificial-intelligence","chatbot","computer-vision","generative-ai","machine-learning","nlp"],"archived":false,"github_pushed_at":"2026-08-24T20:46:59+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/aurelio-labs-semantic-router","markdown_url":"https://www.graphcanon.com/tools/aurelio-labs-semantic-router.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/aurelio-labs-semantic-router","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=aurelio-labs-semantic-router","shared_categories":["llm-frameworks","inference-serving"]},{"slug":"b4rtaz-distributed-llama","name":"distributed-llama","tagline":"Distributed LLM inference using home devices cluster","github_url":"https://github.com/b4rtaz/distributed-llama","owner":"b4rtaz","repo":"distributed-llama","owner_avatar_url":"https://avatars.githubusercontent.com/u/12797776?v=4","primary_language":"C++","stars":3060,"forks":250,"topics":["distributed-computing","distributed-llm","llama2","llama3","llm","llm-inference","llms","neural-network","open-llm"],"archived":false,"github_pushed_at":"2026-07-05T16:47:20+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/b4rtaz-distributed-llama","markdown_url":"https://www.graphcanon.com/tools/b4rtaz-distributed-llama.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/b4rtaz-distributed-llama","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=b4rtaz-distributed-llama","shared_categories":["inference-serving"]},{"slug":"jhubi1-ollama-app","name":"ollama-app","tagline":"A modern and easy-to-use client for Ollama","github_url":"https://github.com/JHubi1/ollama-app","owner":"JHubi1","repo":"ollama-app","owner_avatar_url":"https://avatars.githubusercontent.com/u/61345690?v=4","primary_language":"Dart","stars":1809,"forks":220,"topics":["ai","android","app","linux","llama","localai","ollama","ollama-client","windows"],"archived":false,"github_pushed_at":"2026-08-07T21:26:25+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/jhubi1-ollama-app","markdown_url":"https://www.graphcanon.com/tools/jhubi1-ollama-app.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/jhubi1-ollama-app","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=jhubi1-ollama-app","shared_categories":["llm-frameworks"]},{"slug":"theopenco-llmgateway","name":"llmgateway","tagline":"Route and manage LLM requests via unified API interface","github_url":"https://github.com/theopenco/llmgateway","owner":"theopenco","repo":"llmgateway","owner_avatar_url":"https://avatars.githubusercontent.com/u/211671860?v=4","primary_language":"TypeScript","stars":1624,"forks":181,"topics":["ai","ai-gateway","analytics","anthropic","api-key-management","claude","codex","enterprise","guardrails","inference","llm","llm-gateway","llm-proxy","llms","observability","openai","opencode","rate-limiting","typescript"],"archived":false,"github_pushed_at":"2026-09-11T06:00:04+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/theopenco-llmgateway","markdown_url":"https://www.graphcanon.com/tools/theopenco-llmgateway.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/theopenco-llmgateway","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=theopenco-llmgateway","shared_categories":["inference-serving"]},{"slug":"sauravpanda-browserai","name":"BrowserAI","tagline":"Run local LLMs like llama, deepseek-distill, kokoro and more inside your browser","github_url":"https://github.com/sauravpanda/BrowserAI","owner":"sauravpanda","repo":"BrowserAI","owner_avatar_url":"https://avatars.githubusercontent.com/u/12201824?v=4","primary_language":"TypeScript","stars":1451,"forks":136,"topics":["agents","ai","llama","llm","llm-inference","local","localllm","tts","webgpu"],"archived":false,"github_pushed_at":"2026-07-21T02:47:31+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/sauravpanda-browserai","markdown_url":"https://www.graphcanon.com/tools/sauravpanda-browserai.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/sauravpanda-browserai","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=sauravpanda-browserai","shared_categories":["llm-frameworks","inference-serving"]}]}}