{"data":{"node":{"slug":"notpunchnox-rkllama","name":"rkllama","tagline":"Ollama alternative for Rockchip NPU with optimized AI and Deep learning model inference","github_url":"https://github.com/NotPunchnox/rkllama","owner":"NotPunchnox","repo":"rkllama","owner_avatar_url":"https://avatars.githubusercontent.com/u/74469788?v=4","primary_language":"Python","stars":590,"forks":99,"topics":["ai","client","client-server","ia","llm","llm-apps","llm-inference","npu","npu-llm","offline","orange-pi","orangepi","orangepi5","orangepi5pro","python","rk3576","rk3588","rockchip","server"],"archived":false,"github_pushed_at":"2026-07-07T06:49:18+00:00","maintenance_label":"Steady","stars_delta_30d":13,"url":"https://www.graphcanon.com/tools/notpunchnox-rkllama","markdown_url":"https://www.graphcanon.com/tools/notpunchnox-rkllama.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/notpunchnox-rkllama","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=notpunchnox-rkllama"},"categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"}],"tags":[{"slug":"ai","name":"ai"},{"slug":"client-server","name":"client-server"},{"slug":"llm-inference","name":"llm-inference"},{"slug":"npu-llm","name":"npu-llm"},{"slug":"orange-pi","name":"orange-pi"},{"slug":"orangepi5pro","name":"orangepi5pro"},{"slug":"python","name":"python"},{"slug":"rk3576","name":"rk3576"}],"edges":[],"neighbours":[{"slug":"alexsjones-llmfit","name":"llmfit","tagline":"Hundreds of models & providers. One command to find what runs on your hardware.","github_url":"https://github.com/AlexsJones/llmfit","owner":"AlexsJones","repo":"llmfit","owner_avatar_url":"https://avatars.githubusercontent.com/u/1235925?v=4","primary_language":"Rust","stars":31867,"forks":1978,"topics":["gguf","llm","localai","mlx","skill","unsloth"],"archived":false,"github_pushed_at":"2026-08-14T07:36:41+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/alexsjones-llmfit","markdown_url":"https://www.graphcanon.com/tools/alexsjones-llmfit.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/alexsjones-llmfit","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=alexsjones-llmfit","shared_categories":[]},{"slug":"oumi-ai-oumi","name":"oumi","tagline":"Easily fine-tune, evaluate and deploy open source LLMs/VLMs","github_url":"https://github.com/oumi-ai/oumi","owner":"oumi-ai","repo":"oumi","owner_avatar_url":"https://avatars.githubusercontent.com/u/167452922?v=4","primary_language":"Python","stars":9376,"forks":784,"topics":["dpo","evaluation","fine-tuning","gpt-oss","gpt-oss-120b","gpt-oss-20b","inference","llama","llms","open-weight","open-weight-models","open-weights","sft","slms","vlms"],"archived":false,"github_pushed_at":"2026-08-21T23:11:35+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/oumi-ai-oumi","markdown_url":"https://www.graphcanon.com/tools/oumi-ai-oumi.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/oumi-ai-oumi","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=oumi-ai-oumi","shared_categories":["inference-serving"]},{"slug":"andyyyy64-whichllm","name":"whichllm","tagline":"Command-line tool to find and benchmark local LLM performance","github_url":"https://github.com/Andyyyy64/whichllm","owner":"Andyyyy64","repo":"whichllm","owner_avatar_url":"https://avatars.githubusercontent.com/u/105579829?v=4","primary_language":"Python","stars":6225,"forks":330,"topics":["ai","apple-silicon","benchmarks","cli","command-line-tool","gguf","gpu","huggingface","inference","llm","local-llm","ollama","python","vram"],"archived":false,"github_pushed_at":"2026-08-05T07:15:32+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/andyyyy64-whichllm","markdown_url":"https://www.graphcanon.com/tools/andyyyy64-whichllm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/andyyyy64-whichllm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=andyyyy64-whichllm","shared_categories":["inference-serving"]},{"slug":"b4rtaz-distributed-llama","name":"distributed-llama","tagline":"Distributed LLM inference using home devices cluster","github_url":"https://github.com/b4rtaz/distributed-llama","owner":"b4rtaz","repo":"distributed-llama","owner_avatar_url":"https://avatars.githubusercontent.com/u/12797776?v=4","primary_language":"C++","stars":3044,"forks":246,"topics":["distributed-computing","distributed-llm","llama2","llama3","llm","llm-inference","llms","neural-network","open-llm"],"archived":false,"github_pushed_at":"2026-07-05T16:47:20+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/b4rtaz-distributed-llama","markdown_url":"https://www.graphcanon.com/tools/b4rtaz-distributed-llama.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/b4rtaz-distributed-llama","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=b4rtaz-distributed-llama","shared_categories":["inference-serving"]},{"slug":"off-grid-ai-off-grid-ai-mobile","name":"off-grid-ai-mobile","tagline":"The Swiss Army Knife of Offline AI","github_url":"https://github.com/off-grid-ai/off-grid-ai-mobile","owner":"off-grid-ai","repo":"off-grid-ai-mobile","owner_avatar_url":"https://avatars.githubusercontent.com/u/295106271?v=4","primary_language":"TypeScript","stars":2855,"forks":273,"topics":["android","edge-ai","gguf","ios","llama-cpp","local-ai","macos","mcp","mobile-ai","offline-ai","offline-llm","ondevice","ondevice-ai","privacy-first","react-native","stable-diffusion-android","tool-calling","vision-language-model","whisper-android","whisper-cpp"],"archived":false,"github_pushed_at":"2026-08-01T11:32:40+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/off-grid-ai-off-grid-ai-mobile","markdown_url":"https://www.graphcanon.com/tools/off-grid-ai-off-grid-ai-mobile.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/off-grid-ai-off-grid-ai-mobile","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=off-grid-ai-off-grid-ai-mobile","shared_categories":["inference-serving"]},{"slug":"superlinked-sie","name":"sie","tagline":"Open-source inference server and production cluster for all the models your agent needs.","github_url":"https://github.com/superlinked/sie","owner":"superlinked","repo":"sie","owner_avatar_url":"https://avatars.githubusercontent.com/u/94243920?v=4","primary_language":"Python","stars":2804,"forks":272,"topics":["bge","colbert","data-pipeline","deep-learning","embeddings","inference","inference-server","information-retrieval","llm","ml","mlops","natural-language-processing","nlp","python","reranking","retrieval","retrieval-augmented-generation","semantic-search","splade","vector-search"],"archived":false,"github_pushed_at":"2026-08-21T20:28:04+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/superlinked-sie","markdown_url":"https://www.graphcanon.com/tools/superlinked-sie.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/superlinked-sie","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=superlinked-sie","shared_categories":["inference-serving"]},{"slug":"rafska-awesome-local-llm","name":"awesome-local-llm","tagline":"Resources for running LLMs locally","github_url":"https://github.com/rafska/awesome-local-llm","owner":"rafska","repo":"awesome-local-llm","owner_avatar_url":"https://avatars.githubusercontent.com/u/17859377?v=4","primary_language":null,"stars":2518,"forks":316,"topics":["ai","awesome","awesome-list","llm","local","local-ai","local-llm","resources","self-hosted","selfhosted"],"archived":false,"github_pushed_at":"2026-08-04T23:08:15+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/rafska-awesome-local-llm","markdown_url":"https://www.graphcanon.com/tools/rafska-awesome-local-llm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/rafska-awesome-local-llm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=rafska-awesome-local-llm","shared_categories":["inference-serving"]},{"slug":"jhubi1-ollama-app","name":"ollama-app","tagline":"A modern and easy-to-use client for Ollama","github_url":"https://github.com/JHubi1/ollama-app","owner":"JHubi1","repo":"ollama-app","owner_avatar_url":"https://avatars.githubusercontent.com/u/61345690?v=4","primary_language":"Dart","stars":1797,"forks":216,"topics":["ai","android","app","linux","llama","localai","ollama","ollama-client","windows"],"archived":false,"github_pushed_at":"2026-08-07T21:26:25+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/jhubi1-ollama-app","markdown_url":"https://www.graphcanon.com/tools/jhubi1-ollama-app.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/jhubi1-ollama-app","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=jhubi1-ollama-app","shared_categories":[]},{"slug":"waybarrios-vllm-mlx","name":"vllm-mlx","tagline":"Server for LLMs and vision-language models compatible with Apple Silicon","github_url":"https://github.com/waybarrios/vllm-mlx","owner":"waybarrios","repo":"vllm-mlx","owner_avatar_url":"https://avatars.githubusercontent.com/u/6794828?v=4","primary_language":"Python","stars":1472,"forks":205,"topics":["anthropic","apple-silicon","audio-processing","claude-code","computer-vision","image-understanding","inference","llm","machine-learning","macos","mllm","mlx","multimodal-ai","speech-to-text","stt","text-to-speech","tts","video-understanding","vision-language-model","vllm"],"archived":false,"github_pushed_at":"2026-06-28T20:18:31+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/waybarrios-vllm-mlx","markdown_url":"https://www.graphcanon.com/tools/waybarrios-vllm-mlx.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/waybarrios-vllm-mlx","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=waybarrios-vllm-mlx","shared_categories":["inference-serving"]},{"slug":"sauravpanda-browserai","name":"BrowserAI","tagline":"Run local LLMs like llama, deepseek-distill, kokoro and more inside your browser","github_url":"https://github.com/sauravpanda/BrowserAI","owner":"sauravpanda","repo":"BrowserAI","owner_avatar_url":"https://avatars.githubusercontent.com/u/12201824?v=4","primary_language":"TypeScript","stars":1449,"forks":138,"topics":["agents","ai","llama","llm","llm-inference","local","localllm","tts","webgpu"],"archived":false,"github_pushed_at":"2026-07-21T02:47:31+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/sauravpanda-browserai","markdown_url":"https://www.graphcanon.com/tools/sauravpanda-browserai.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/sauravpanda-browserai","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=sauravpanda-browserai","shared_categories":["inference-serving"]},{"slug":"kubeai-project-kubeai","name":"kubeai","tagline":"AI Inference Operator for Kubernetes","github_url":"https://github.com/kubeai-project/kubeai","owner":"kubeai-project","repo":"kubeai","owner_avatar_url":"https://avatars.githubusercontent.com/u/232319222?v=4","primary_language":"Go","stars":1237,"forks":131,"topics":["ai","autoscaler","faster-whisper","inference-operator","k8s","kubernetes","llm","ollama","ollama-operator","openai-api","vllm","vllm-operator","whisper"],"archived":false,"github_pushed_at":"2026-07-31T01:04:47+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/kubeai-project-kubeai","markdown_url":"https://www.graphcanon.com/tools/kubeai-project-kubeai.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/kubeai-project-kubeai","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=kubeai-project-kubeai","shared_categories":["inference-serving"]},{"slug":"jonigl-mcp-client-for-ollama","name":"mcp-client-for-ollama","tagline":"TUI MCP Client for Ollama enables local LLM interaction with extensive features.","github_url":"https://github.com/jonigl/mcp-client-for-ollama","owner":"jonigl","repo":"mcp-client-for-ollama","owner_avatar_url":"https://avatars.githubusercontent.com/u/4612832?v=4","primary_language":"Python","stars":783,"forks":114,"topics":["agentic-ai","ai","command-line-tool","harness","linux","local-llm","macos","mcp","mcp-client","mcp-prompts","mcp-resouces","mcp-server","mcp-tools","model-context-protocol","ollama","open-source","sse","stdio","streamable-http","windows"],"archived":false,"github_pushed_at":"2026-07-27T08:34:02+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/jonigl-mcp-client-for-ollama","markdown_url":"https://www.graphcanon.com/tools/jonigl-mcp-client-for-ollama.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/jonigl-mcp-client-for-ollama","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=jonigl-mcp-client-for-ollama","shared_categories":["inference-serving"]}]}}