{"data":{"node":{"slug":"techopolis-afm-server","name":"afm-Server","tagline":"macOS menu bar app for exposing Apple's on-device Foundation Models via an OpenAI-compatible API","github_url":"https://github.com/Techopolis/afm-Server","owner":"Techopolis","repo":"afm-Server","owner_avatar_url":"https://avatars.githubusercontent.com/u/71575185?v=4","primary_language":"Swift","stars":192,"forks":9,"topics":["apple-intelligence","foundation-models","local-llm","macos","menu-bar-app","on-device-ai","openai-api","swift"],"archived":false,"github_pushed_at":"2026-06-02T03:48:27+00:00","maintenance_label":"Slowing","stars_delta_30d":3,"url":"https://www.graphcanon.com/tools/techopolis-afm-server","markdown_url":"https://www.graphcanon.com/tools/techopolis-afm-server.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/techopolis-afm-server","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=techopolis-afm-server"},"categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"}],"tags":[{"slug":"apple-intelligence","name":"apple-intelligence"},{"slug":"foundation-models","name":"foundation-models"},{"slug":"local-llm","name":"local-llm"},{"slug":"macos","name":"macos"},{"slug":"menu-bar-app","name":"menu-bar-app"},{"slug":"on-device-ai","name":"on-device-ai"},{"slug":"openai-api","name":"openai-api"},{"slug":"swift","name":"swift"}],"edges":[],"neighbours":[{"slug":"open-webui-open-webui","name":"open-webui","tagline":"User-friendly AI Interface","github_url":"https://github.com/open-webui/open-webui","owner":"open-webui","repo":"open-webui","owner_avatar_url":"https://avatars.githubusercontent.com/u/158137808?v=4","primary_language":"Python","stars":152445,"forks":22306,"topics":["ai","llm","llm-ui","llm-webui","llms","mcp","ollama","ollama-webui","open-webui","openai","openapi","rag","self-hosted","ui","webui"],"archived":false,"github_pushed_at":"2026-09-18T00:11:06+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/open-webui-open-webui","markdown_url":"https://www.graphcanon.com/tools/open-webui-open-webui.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/open-webui-open-webui","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=open-webui-open-webui","shared_categories":["inference-serving"]},{"slug":"jundot-omlx","name":"omlx","tagline":"LLM inference server with continuous batching and SSD caching for Apple Silicon","github_url":"https://github.com/jundot/omlx","owner":"jundot","repo":"omlx","owner_avatar_url":"https://avatars.githubusercontent.com/u/64250138?v=4","primary_language":"Python","stars":21934,"forks":1899,"topics":["apple-silicon","inference-server","llm","macos","mlx","openai-api"],"archived":false,"github_pushed_at":"2026-09-20T02:23:26+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/jundot-omlx","markdown_url":"https://www.graphcanon.com/tools/jundot-omlx.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/jundot-omlx","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=jundot-omlx","shared_categories":["inference-serving"]},{"slug":"superlinked-sie","name":"sie","tagline":"Open-source inference server and production cluster for all the models your agent needs.","github_url":"https://github.com/superlinked/sie","owner":"superlinked","repo":"sie","owner_avatar_url":"https://avatars.githubusercontent.com/u/94243920?v=4","primary_language":"Python","stars":3293,"forks":311,"topics":["bge","colbert","data-pipeline","deep-learning","embeddings","inference","inference-server","information-retrieval","llm","ml","mlops","natural-language-processing","nlp","python","reranking","retrieval","retrieval-augmented-generation","semantic-search","splade","vector-search"],"archived":false,"github_pushed_at":"2026-09-19T01:00:02+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/superlinked-sie","markdown_url":"https://www.graphcanon.com/tools/superlinked-sie.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/superlinked-sie","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=superlinked-sie","shared_categories":["inference-serving"]},{"slug":"off-grid-ai-off-grid-ai-mobile","name":"OGAM","tagline":"Swiss Army Knife of Offline AI","github_url":"https://github.com/off-grid-ai/OGAM","owner":"off-grid-ai","repo":"OGAM","owner_avatar_url":"https://avatars.githubusercontent.com/u/295106271?v=4","primary_language":"TypeScript","stars":3124,"forks":298,"topics":["android","edge-ai","gguf","ios","llama-cpp","local-ai","macos","mcp","mobile-ai","offline-ai","offline-llm","ondevice","ondevice-ai","privacy-first","react-native","stable-diffusion-android","tool-calling","vision-language-model","whisper-android","whisper-cpp"],"archived":false,"github_pushed_at":"2026-09-18T07:38:15+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/off-grid-ai-off-grid-ai-mobile","markdown_url":"https://www.graphcanon.com/tools/off-grid-ai-off-grid-ai-mobile.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/off-grid-ai-off-grid-ai-mobile","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=off-grid-ai-off-grid-ai-mobile","shared_categories":["inference-serving"]},{"slug":"b4rtaz-distributed-llama","name":"distributed-llama","tagline":"Distributed LLM inference using home devices cluster","github_url":"https://github.com/b4rtaz/distributed-llama","owner":"b4rtaz","repo":"distributed-llama","owner_avatar_url":"https://avatars.githubusercontent.com/u/12797776?v=4","primary_language":"C++","stars":3060,"forks":250,"topics":["distributed-computing","distributed-llm","llama2","llama3","llm","llm-inference","llms","neural-network","open-llm"],"archived":false,"github_pushed_at":"2026-07-05T16:47:20+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/b4rtaz-distributed-llama","markdown_url":"https://www.graphcanon.com/tools/b4rtaz-distributed-llama.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/b4rtaz-distributed-llama","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=b4rtaz-distributed-llama","shared_categories":["inference-serving"]},{"slug":"atomicbot-ai-atomic-chat","name":"Atomic-Chat","tagline":"Local AI app and inference engine for agents","github_url":"https://github.com/AtomicBot-ai/Atomic-Chat","owner":"AtomicBot-ai","repo":"Atomic-Chat","owner_avatar_url":"https://avatars.githubusercontent.com/u/259533419?v=4","primary_language":"TypeScript","stars":1526,"forks":178,"topics":["ai-chat","ai-tools","apple-silicon","chatgpt","deepseek","desktop-app","gemma","gguf","gpt-oss","llamacpp","llm","llm-inference","local-ai","local-first","local-llm","mcp","mlx","open-source","qwen","self-hosted"],"archived":false,"github_pushed_at":"2026-09-19T07:46:33+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/atomicbot-ai-atomic-chat","markdown_url":"https://www.graphcanon.com/tools/atomicbot-ai-atomic-chat.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/atomicbot-ai-atomic-chat","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=atomicbot-ai-atomic-chat","shared_categories":["inference-serving"]},{"slug":"ddalcu-mlx-serve","name":"mlx-serve","tagline":"Native LLM inference server for Apple Silicon","github_url":"https://github.com/ddalcu/mlx-serve","owner":"ddalcu","repo":"mlx-serve","owner_avatar_url":"https://avatars.githubusercontent.com/u/869085?v=4","primary_language":"Zig","stars":1418,"forks":130,"topics":["agent","anthropic-api","apple-silicon","claude-code","deepseek-v4","diffusion","gguf","image-generation","inference","llm","local-llm","macos","macos-app","mlx","openai-api","tool-calling","video-generation","voice-agent","voice-cloning","zig"],"archived":false,"github_pushed_at":"2026-09-19T22:08:35+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/ddalcu-mlx-serve","markdown_url":"https://www.graphcanon.com/tools/ddalcu-mlx-serve.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ddalcu-mlx-serve","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ddalcu-mlx-serve","shared_categories":["inference-serving"]},{"slug":"rudrankriyam-foundation-models-framework-lab","name":"Foundation-Models-Framework-Lab","tagline":"A practical lab for building, testing, and evaluating apps with Apple's Foundation Models framework","github_url":"https://github.com/rudrankriyam/Foundation-Models-Framework-Lab","owner":"rudrankriyam","repo":"Foundation-Models-Framework-Lab","owner_avatar_url":"https://avatars.githubusercontent.com/u/30552772?v=4","primary_language":"Swift","stars":1181,"forks":70,"topics":["ai","apple-foundation-models","apple-intelligence","foundation-models","foundation-models-framework","generative-ai","healthkit","ios","large-language-models","llm","macos","multilingual","on-device-ai","rag","speech-recognition","swift","swiftui","text-to-speech","tool-calling","xcode"],"archived":false,"github_pushed_at":"2026-09-18T23:34:01+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/rudrankriyam-foundation-models-framework-lab","markdown_url":"https://www.graphcanon.com/tools/rudrankriyam-foundation-models-framework-lab.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/rudrankriyam-foundation-models-framework-lab","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=rudrankriyam-foundation-models-framework-lab","shared_categories":[]},{"slug":"timmyy123-llm-hub","name":"LLM-Hub","tagline":"Local AI Assistant on your phone","github_url":"https://github.com/timmyy123/LLM-Hub","owner":"timmyy123","repo":"LLM-Hub","owner_avatar_url":"https://avatars.githubusercontent.com/u/164847362?v=4","primary_language":"C++","stars":588,"forks":124,"topics":["ai","gemma","gemma4","gemma4-agent-skills","gptoss","granite","imagegeneration","lfm25","llama","llm","llm-inference","mistral","music","musicgeneration","phi4","rag","stable-diffusion","videogeneration","whisper"],"archived":false,"github_pushed_at":"2026-09-19T11:49:49+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/timmyy123-llm-hub","markdown_url":"https://www.graphcanon.com/tools/timmyy123-llm-hub.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/timmyy123-llm-hub","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=timmyy123-llm-hub","shared_categories":["inference-serving"]},{"slug":"kaito-project-aikit","name":"aikit","tagline":"Fine-tune, build, and deploy open-source LLMs easily!","github_url":"https://github.com/kaito-project/aikit","owner":"kaito-project","repo":"aikit","owner_avatar_url":"https://avatars.githubusercontent.com/u/186863079?v=4","primary_language":"Go","stars":539,"forks":57,"topics":["ai","buildkit","chatgpt","docker","fine-tuning","finetuning","gemma","gpt","inference","kubernetes","large-language-models","llama","llm","localllama","mistral","mixtral","nvidia","open-llm","open-source-llm","openai"],"archived":false,"github_pushed_at":"2026-09-18T22:43:05+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/kaito-project-aikit","markdown_url":"https://www.graphcanon.com/tools/kaito-project-aikit.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/kaito-project-aikit","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=kaito-project-aikit","shared_categories":["inference-serving"]},{"slug":"withcatai-catai","name":"catai","tagline":"Run AI assistant locally with Node.js","github_url":"https://github.com/withcatai/catai","owner":"withcatai","repo":"catai","owner_avatar_url":"https://avatars.githubusercontent.com/u/142035548?v=4","primary_language":"TypeScript","stars":500,"forks":39,"topics":["ai","ai-assistant","catai","chatbot","chatgpt","chatui","dalai","ggmlv3","gguf","llama-cpp","llm","local-llm","localai","node-llama-cpp","nodejs","openai","vicuna","vicuna-installation-guide","wizardlm"],"archived":false,"github_pushed_at":"2025-11-16T18:35:45+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/withcatai-catai","markdown_url":"https://www.graphcanon.com/tools/withcatai-catai.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/withcatai-catai","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=withcatai-catai","shared_categories":["inference-serving"]},{"slug":"writesonic-gptrouter","name":"GPTRouter","tagline":"Manage multiple LLMs and image models for reliable and fast responses","github_url":"https://github.com/Writesonic/GPTRouter","owner":"Writesonic","repo":"GPTRouter","owner_avatar_url":"https://avatars.githubusercontent.com/u/83169781?v=4","primary_language":"TypeScript","stars":456,"forks":38,"topics":["anthropic","azure-openai","cohere","google-gemini","langchain","llama-index","llm","llmops","llms","mlops","openai","palm-api"],"archived":false,"github_pushed_at":"2024-04-10T11:14:59+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/writesonic-gptrouter","markdown_url":"https://www.graphcanon.com/tools/writesonic-gptrouter.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/writesonic-gptrouter","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=writesonic-gptrouter","shared_categories":["inference-serving"]}]}}