{"data":{"node":{"slug":"devnen-qwen3-6-windows-server","name":"qwen3.6-windows-server","tagline":"One-click Qwen3.6-27B inference tool for Windows","github_url":"https://github.com/devnen/qwen3.6-windows-server","owner":"devnen","repo":"qwen3.6-windows-server","owner_avatar_url":"https://avatars.githubusercontent.com/u/195903272?v=4","primary_language":"Python","stars":229,"forks":23,"topics":["llm-inference","local-llm","offline-ai","privacy","qwen","qwen3","rtx-3090","textual-tui","vllm","windows"],"archived":false,"github_pushed_at":"2026-05-14T17:14:43+00:00","maintenance_label":"Slowing","stars_delta_30d":2,"url":"https://www.graphcanon.com/tools/devnen-qwen3-6-windows-server","markdown_url":"https://www.graphcanon.com/tools/devnen-qwen3-6-windows-server.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/devnen-qwen3-6-windows-server","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=devnen-qwen3-6-windows-server"},"categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"}],"tags":[{"slug":"llm-inference","name":"llm-inference"},{"slug":"local-llm","name":"local-llm"},{"slug":"offline-ai","name":"offline-ai"},{"slug":"privacy","name":"privacy"},{"slug":"qwen","name":"qwen"},{"slug":"vllm","name":"vllm"},{"slug":"windows","name":"windows"}],"edges":[],"neighbours":[{"slug":"gaizhenbiao-chuanhuchatgpt","name":"ChuanhuChatGPT","tagline":"GUI for ChatGPT and LLMs with features like agents and file-based QA.","github_url":"https://github.com/GaiZhenbiao/ChuanhuChatGPT","owner":"GaiZhenbiao","repo":"ChuanhuChatGPT","owner_avatar_url":"https://avatars.githubusercontent.com/u/51039745?v=4","primary_language":"Python","stars":15272,"forks":2195,"topics":["chatbot","chatglm","chatgpt-api","claude","dalle3","ernie","gemini","gemma","inspurai","llama","midjourney","minimax","moss","ollama","qwen","spark","stablelm"],"archived":false,"github_pushed_at":"2026-09-16T11:44:01+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/gaizhenbiao-chuanhuchatgpt","markdown_url":"https://www.graphcanon.com/tools/gaizhenbiao-chuanhuchatgpt.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/gaizhenbiao-chuanhuchatgpt","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=gaizhenbiao-chuanhuchatgpt","shared_categories":["inference-serving"]},{"slug":"xorbitsai-inference","name":"inference","tagline":"Unified production-ready inference API for various LLMs and models","github_url":"https://github.com/xorbitsai/inference","owner":"xorbitsai","repo":"inference","owner_avatar_url":"https://avatars.githubusercontent.com/u/109655068?v=4","primary_language":"Python","stars":9575,"forks":866,"topics":["artificial-intelligence","deployment","diffusers","gemma","glm","glm-5-3","inference","kimi","kimi-k3","llama-cpp","llamacpp","llm","machine-learning","openai-api","pytorch","qwen","sglang","transformers","vllm","whisper"],"archived":false,"github_pushed_at":"2026-09-18T06:04:24+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/xorbitsai-inference","markdown_url":"https://www.graphcanon.com/tools/xorbitsai-inference.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/xorbitsai-inference","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=xorbitsai-inference","shared_categories":["inference-serving"]},{"slug":"oumi-ai-oumi","name":"oumi","tagline":"Easily fine-tune, evaluate and deploy open source LLMs/VLMs","github_url":"https://github.com/oumi-ai/oumi","owner":"oumi-ai","repo":"oumi","owner_avatar_url":"https://avatars.githubusercontent.com/u/167452922?v=4","primary_language":"Python","stars":9387,"forks":790,"topics":["dpo","evaluation","fine-tuning","gpt-oss","gpt-oss-120b","gpt-oss-20b","inference","llama","llms","open-weight","open-weight-models","open-weights","sft","slms","vlms"],"archived":false,"github_pushed_at":"2026-09-20T05:45:37+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/oumi-ai-oumi","markdown_url":"https://www.graphcanon.com/tools/oumi-ai-oumi.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/oumi-ai-oumi","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=oumi-ai-oumi","shared_categories":["inference-serving"]},{"slug":"andyyyy64-whichllm","name":"whichllm","tagline":"Command-line tool to find and benchmark local LLM performance","github_url":"https://github.com/Andyyyy64/whichllm","owner":"Andyyyy64","repo":"whichllm","owner_avatar_url":"https://avatars.githubusercontent.com/u/105579829?v=4","primary_language":"Python","stars":6666,"forks":368,"topics":["localllm"],"archived":false,"github_pushed_at":"2026-09-19T16:22:49+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/andyyyy64-whichllm","markdown_url":"https://www.graphcanon.com/tools/andyyyy64-whichllm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/andyyyy64-whichllm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=andyyyy64-whichllm","shared_categories":["inference-serving"]},{"slug":"osmantic-ods","name":"ODS","tagline":"Transform personal computers into AI servers.","github_url":"https://github.com/Osmantic/ODS","owner":"Osmantic","repo":"ODS","owner_avatar_url":"https://avatars.githubusercontent.com/u/262014141?v=4","primary_language":"Python","stars":6607,"forks":943,"topics":["ai-agents","amd","comfyui","docker","llama-cpp","llm","local-ai","n8n","nvidia","open-webui","rag","self-hosted","speech-to-text","strix-halo","text-to-speech","workflow-automation"],"archived":false,"github_pushed_at":"2026-09-20T07:02:35+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/osmantic-ods","markdown_url":"https://www.graphcanon.com/tools/osmantic-ods.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/osmantic-ods","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=osmantic-ods","shared_categories":["inference-serving"]},{"slug":"superlinked-sie","name":"sie","tagline":"Open-source inference server and production cluster for all the models your agent needs.","github_url":"https://github.com/superlinked/sie","owner":"superlinked","repo":"sie","owner_avatar_url":"https://avatars.githubusercontent.com/u/94243920?v=4","primary_language":"Python","stars":3293,"forks":311,"topics":["bge","colbert","data-pipeline","deep-learning","embeddings","inference","inference-server","information-retrieval","llm","ml","mlops","natural-language-processing","nlp","python","reranking","retrieval","retrieval-augmented-generation","semantic-search","splade","vector-search"],"archived":false,"github_pushed_at":"2026-09-19T01:00:02+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/superlinked-sie","markdown_url":"https://www.graphcanon.com/tools/superlinked-sie.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/superlinked-sie","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=superlinked-sie","shared_categories":["inference-serving"]},{"slug":"b4rtaz-distributed-llama","name":"distributed-llama","tagline":"Distributed LLM inference using home devices cluster","github_url":"https://github.com/b4rtaz/distributed-llama","owner":"b4rtaz","repo":"distributed-llama","owner_avatar_url":"https://avatars.githubusercontent.com/u/12797776?v=4","primary_language":"C++","stars":3060,"forks":250,"topics":["distributed-computing","distributed-llm","llama2","llama3","llm","llm-inference","llms","neural-network","open-llm"],"archived":false,"github_pushed_at":"2026-07-05T16:47:20+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/b4rtaz-distributed-llama","markdown_url":"https://www.graphcanon.com/tools/b4rtaz-distributed-llama.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/b4rtaz-distributed-llama","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=b4rtaz-distributed-llama","shared_categories":["inference-serving"]},{"slug":"rafska-awesome-local-llm","name":"awesome-local-llm","tagline":"Resources for running LLMs locally","github_url":"https://github.com/rafska/awesome-local-llm","owner":"rafska","repo":"awesome-local-llm","owner_avatar_url":"https://avatars.githubusercontent.com/u/17859377?v=4","primary_language":null,"stars":2869,"forks":388,"topics":["ai","awesome","awesome-list","llm","local","local-ai","local-llm","resources","self-hosted","selfhosted"],"archived":false,"github_pushed_at":"2026-09-13T09:48:26+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/rafska-awesome-local-llm","markdown_url":"https://www.graphcanon.com/tools/rafska-awesome-local-llm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/rafska-awesome-local-llm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=rafska-awesome-local-llm","shared_categories":["inference-serving"]},{"slug":"waybarrios-vllm-mlx","name":"vllm-mlx","tagline":"Server for LLMs and vision-language models compatible with Apple Silicon","github_url":"https://github.com/waybarrios/vllm-mlx","owner":"waybarrios","repo":"vllm-mlx","owner_avatar_url":"https://avatars.githubusercontent.com/u/6794828?v=4","primary_language":"Python","stars":1588,"forks":222,"topics":["anthropic","anthropic-api","apple-silicon","claude-code","continuous-batching","inference-server","llm","local-llm","macos","mcp","mlx","multimodal-ai","openai","openai-api","openai-compatible","speech-to-text","text-to-speech","tool-calling","vision-language-model","vllm"],"archived":false,"github_pushed_at":"2026-09-19T18:11:07+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/waybarrios-vllm-mlx","markdown_url":"https://www.graphcanon.com/tools/waybarrios-vllm-mlx.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/waybarrios-vllm-mlx","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=waybarrios-vllm-mlx","shared_categories":["inference-serving"]},{"slug":"atomicbot-ai-atomic-chat","name":"Atomic-Chat","tagline":"Local AI app and inference engine for agents","github_url":"https://github.com/AtomicBot-ai/Atomic-Chat","owner":"AtomicBot-ai","repo":"Atomic-Chat","owner_avatar_url":"https://avatars.githubusercontent.com/u/259533419?v=4","primary_language":"TypeScript","stars":1526,"forks":178,"topics":["ai-chat","ai-tools","apple-silicon","chatgpt","deepseek","desktop-app","gemma","gguf","gpt-oss","llamacpp","llm","llm-inference","local-ai","local-first","local-llm","mcp","mlx","open-source","qwen","self-hosted"],"archived":false,"github_pushed_at":"2026-09-19T07:46:33+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/atomicbot-ai-atomic-chat","markdown_url":"https://www.graphcanon.com/tools/atomicbot-ai-atomic-chat.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/atomicbot-ai-atomic-chat","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=atomicbot-ai-atomic-chat","shared_categories":["inference-serving"]},{"slug":"sauravpanda-browserai","name":"BrowserAI","tagline":"Run local LLMs like llama, deepseek-distill, kokoro and more inside your browser","github_url":"https://github.com/sauravpanda/BrowserAI","owner":"sauravpanda","repo":"BrowserAI","owner_avatar_url":"https://avatars.githubusercontent.com/u/12201824?v=4","primary_language":"TypeScript","stars":1451,"forks":136,"topics":["agents","ai","llama","llm","llm-inference","local","localllm","tts","webgpu"],"archived":false,"github_pushed_at":"2026-07-21T02:47:31+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/sauravpanda-browserai","markdown_url":"https://www.graphcanon.com/tools/sauravpanda-browserai.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/sauravpanda-browserai","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=sauravpanda-browserai","shared_categories":["inference-serving"]},{"slug":"ddalcu-mlx-serve","name":"mlx-serve","tagline":"Native LLM inference server for Apple Silicon","github_url":"https://github.com/ddalcu/mlx-serve","owner":"ddalcu","repo":"mlx-serve","owner_avatar_url":"https://avatars.githubusercontent.com/u/869085?v=4","primary_language":"Zig","stars":1418,"forks":130,"topics":["agent","anthropic-api","apple-silicon","claude-code","deepseek-v4","diffusion","gguf","image-generation","inference","llm","local-llm","macos","macos-app","mlx","openai-api","tool-calling","video-generation","voice-agent","voice-cloning","zig"],"archived":false,"github_pushed_at":"2026-09-19T22:08:35+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/ddalcu-mlx-serve","markdown_url":"https://www.graphcanon.com/tools/ddalcu-mlx-serve.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ddalcu-mlx-serve","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ddalcu-mlx-serve","shared_categories":["inference-serving"]}]}}