{"data":{"node":{"slug":"cactus-compute-cactus","name":"cactus","tagline":"Low-latency AI engine for mobile devices & wearables","github_url":"https://github.com/cactus-compute/cactus","owner":"cactus-compute","repo":"cactus","owner_avatar_url":"https://avatars.githubusercontent.com/u/196640840?v=4","primary_language":"C++","stars":5909,"forks":493,"topics":["ai","android","arm","edge","edge-ai","framework","ios","llamacpp","llm","llm-inference","llms","mobile","mobile-inference","on-device-ai","quantiz","rag","smartphone","speech","transformer","whisper"],"archived":false,"github_pushed_at":"2026-08-24T16:05:20+00:00","maintenance_label":"Very active","stars_delta_30d":374,"url":"https://www.graphcanon.com/tools/cactus-compute-cactus","markdown_url":"https://www.graphcanon.com/tools/cactus-compute-cactus.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/cactus-compute-cactus","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=cactus-compute-cactus"},"categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"},{"slug":"speech-audio","name":"Speech & Audio","url":"https://www.graphcanon.com/categories/speech-audio","markdown_url":"https://www.graphcanon.com/categories/speech-audio.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/speech-audio"}],"tags":[{"slug":"ai","name":"ai"},{"slug":"android","name":"android"},{"slug":"arm","name":"arm"},{"slug":"edge","name":"edge"},{"slug":"edge-ai","name":"edge-ai"},{"slug":"framework","name":"framework"},{"slug":"ios","name":"ios"},{"slug":"llamacpp","name":"llamacpp"}],"edges":[],"neighbours":[{"slug":"ollama-ollama","name":"ollama","tagline":"Get up and running with various large language models using Ollama.","github_url":"https://github.com/ollama/ollama","owner":"ollama","repo":"ollama","owner_avatar_url":"https://avatars.githubusercontent.com/u/151674099?v=4","primary_language":"Go","stars":177524,"forks":17229,"topics":["deepseek","gemma","gemma3","glm","go","golang","gpt-oss","llama","llama3","llm","llms","minimax","mistral","ollama","qwen"],"archived":false,"github_pushed_at":"2026-07-31T23:59:29+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/ollama-ollama","markdown_url":"https://www.graphcanon.com/tools/ollama-ollama.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ollama-ollama","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ollama-ollama","shared_categories":["inference-serving"]},{"slug":"huggingface-transformers","name":"transformers","tagline":"Transformers: the model-definition framework for state-of-the-art machine learning models in text, vision, audio, and multimodal models","github_url":"https://github.com/huggingface/transformers","owner":"huggingface","repo":"transformers","owner_avatar_url":"https://avatars.githubusercontent.com/u/25720743?v=4","primary_language":"Python","stars":164121,"forks":34249,"topics":["audio","deep-learning","deepseek","gemma","glm","hacktoberfest","llm","machine-learning","model-hub","natural-language-processing","nlp","pretrained-models","python","pytorch","pytorch-transformers","qwen","speech-recognition","transformer","vlm"],"archived":false,"github_pushed_at":"2026-08-15T22:28:12+00:00","maintenance_label":"Very active","stars_delta_30d":1457,"url":"https://www.graphcanon.com/tools/huggingface-transformers","markdown_url":"https://www.graphcanon.com/tools/huggingface-transformers.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/huggingface-transformers","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=huggingface-transformers","shared_categories":["inference-serving","speech-audio"]},{"slug":"open-webui-open-webui","name":"open-webui","tagline":"User-friendly AI Interface (Supports Ollama, OpenAI API, ...)","github_url":"https://github.com/open-webui/open-webui","owner":"open-webui","repo":"open-webui","owner_avatar_url":"https://avatars.githubusercontent.com/u/158137808?v=4","primary_language":"Python","stars":148875,"forks":21676,"topics":["ai","llm","llm-ui","llm-webui","llms","mcp","ollama","ollama-webui","open-webui","openai","openapi","rag","self-hosted","ui","webui"],"archived":false,"github_pushed_at":"2026-08-15T07:10:16+00:00","maintenance_label":"Very active","stars_delta_30d":3224,"url":"https://www.graphcanon.com/tools/open-webui-open-webui","markdown_url":"https://www.graphcanon.com/tools/open-webui-open-webui.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/open-webui-open-webui","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=open-webui-open-webui","shared_categories":["inference-serving"]},{"slug":"ggml-org-llama-cpp","name":"llama.cpp","tagline":"LLM inference in C/C++","github_url":"https://github.com/ggml-org/llama.cpp","owner":"ggml-org","repo":"llama.cpp","owner_avatar_url":"https://avatars.githubusercontent.com/u/134263123?v=4","primary_language":"C++","stars":122941,"forks":21406,"topics":["ggml"],"archived":false,"github_pushed_at":"2026-08-07T05:28:54+00:00","maintenance_label":"Very active","stars_delta_30d":3353,"url":"https://www.graphcanon.com/tools/ggml-org-llama-cpp","markdown_url":"https://www.graphcanon.com/tools/ggml-org-llama-cpp.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ggml-org-llama-cpp","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ggml-org-llama-cpp","shared_categories":["inference-serving"]},{"slug":"openai-whisper","name":"whisper","tagline":"Robust Speech Recognition via Large-Scale Weak Supervision","github_url":"https://github.com/openai/whisper","owner":"openai","repo":"whisper","owner_avatar_url":"https://avatars.githubusercontent.com/u/14957082?v=4","primary_language":"Python","stars":106740,"forks":12971,"topics":[],"archived":false,"github_pushed_at":"2026-07-28T20:18:29+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/openai-whisper","markdown_url":"https://www.graphcanon.com/tools/openai-whisper.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/openai-whisper","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=openai-whisper","shared_categories":["speech-audio"]},{"slug":"deepseek-ai-deepseek-v3","name":"DeepSeek-V3","tagline":"Repository lacking description with unspecified content related to AI development.","github_url":"https://github.com/deepseek-ai/DeepSeek-V3","owner":"deepseek-ai","repo":"DeepSeek-V3","owner_avatar_url":"https://avatars.githubusercontent.com/u/148330874?v=4","primary_language":"Python","stars":104121,"forks":16726,"topics":[],"archived":false,"github_pushed_at":"2025-08-28T03:24:37+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/deepseek-ai-deepseek-v3","markdown_url":"https://www.graphcanon.com/tools/deepseek-ai-deepseek-v3.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/deepseek-ai-deepseek-v3","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=deepseek-ai-deepseek-v3","shared_categories":["inference-serving"]},{"slug":"pytorch-pytorch","name":"pytorch","tagline":"Tensors and Dynamic neural networks in Python with strong GPU acceleration","github_url":"https://github.com/pytorch/pytorch","owner":"pytorch","repo":"pytorch","owner_avatar_url":"https://avatars.githubusercontent.com/u/21003710?v=4","primary_language":"Python","stars":102144,"forks":28650,"topics":["autograd","deep-learning","gpu","machine-learning","neural-network","numpy","python","tensor"],"archived":false,"github_pushed_at":"2026-08-03T06:00:50+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/pytorch-pytorch","markdown_url":"https://www.graphcanon.com/tools/pytorch-pytorch.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/pytorch-pytorch","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=pytorch-pytorch","shared_categories":["inference-serving"]},{"slug":"thedotmack-claude-mem","name":"claude-mem","tagline":"Persistent Context Across Sessions for Every Agent","github_url":"https://github.com/thedotmack/claude-mem","owner":"thedotmack","repo":"claude-mem","owner_avatar_url":"https://avatars.githubusercontent.com/u/683968?v=4","primary_language":"JavaScript","stars":91018,"forks":7953,"topics":["ai","ai-agents","ai-memory","anthropic","artificial-intelligence","chromadb","claude","claude-agent-sdk","claude-agents","claude-code","claude-code-plugin","claude-skills","embeddings","long-term-memory","mem0","memory-engine","openmemory","rag","sqlite","supermemory"],"archived":false,"github_pushed_at":"2026-08-17T15:46:17+00:00","maintenance_label":"Very active","stars_delta_30d":3316,"url":"https://www.graphcanon.com/tools/thedotmack-claude-mem","markdown_url":"https://www.graphcanon.com/tools/thedotmack-claude-mem.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/thedotmack-claude-mem","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=thedotmack-claude-mem","shared_categories":["inference-serving"]},{"slug":"vllm-project-vllm","name":"vllm","tagline":"A high-throughput and memory-efficient inference and serving engine for LLMs","github_url":"https://github.com/vllm-project/vllm","owner":"vllm-project","repo":"vllm","owner_avatar_url":"https://avatars.githubusercontent.com/u/136984999?v=4","primary_language":"Python","stars":87847,"forks":20135,"topics":["amd","blackwell","cuda","deepseek","deepseek-v3","gpt","gpt-oss","inference","kimi","llama","llm","llm-serving","model-serving","moe","openai","pytorch","qwen","qwen3","tpu","transformer"],"archived":false,"github_pushed_at":"2026-08-01T11:55:36+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/vllm-project-vllm","markdown_url":"https://www.graphcanon.com/tools/vllm-project-vllm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/vllm-project-vllm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=vllm-project-vllm","shared_categories":["inference-serving"]},{"slug":"mlabonne-llm-course","name":"llm-course","tagline":"Course to get into Large Language Models (LLMs) with roadmaps and Colab notebooks.","github_url":"https://github.com/mlabonne/llm-course","owner":"mlabonne","repo":"llm-course","owner_avatar_url":"https://avatars.githubusercontent.com/u/81252890?v=4","primary_language":null,"stars":81512,"forks":9490,"topics":["course","large-language-models","llm","machine-learning","roadmap"],"archived":false,"github_pushed_at":"2026-02-05T13:09:26+00:00","maintenance_label":"Slowing","stars_delta_30d":771,"url":"https://www.graphcanon.com/tools/mlabonne-llm-course","markdown_url":"https://www.graphcanon.com/tools/mlabonne-llm-course.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/mlabonne-llm-course","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=mlabonne-llm-course","shared_categories":["inference-serving"]},{"slug":"nomic-ai-gpt4all","name":"gpt4all","tagline":"Run Local LLMs on Any Device","github_url":"https://github.com/nomic-ai/gpt4all","owner":"nomic-ai","repo":"gpt4all","owner_avatar_url":"https://avatars.githubusercontent.com/u/102670180?v=4","primary_language":"C++","stars":77393,"forks":8296,"topics":["ai-chat","llm-inference"],"archived":false,"github_pushed_at":"2025-05-27T20:05:19+00:00","maintenance_label":"Dormant","stars_delta_30d":-3,"url":"https://www.graphcanon.com/tools/nomic-ai-gpt4all","markdown_url":"https://www.graphcanon.com/tools/nomic-ai-gpt4all.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/nomic-ai-gpt4all","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=nomic-ai-gpt4all","shared_categories":["inference-serving"]},{"slug":"unslothai-unsloth","name":"unsloth","tagline":"A web UI for training and running open models locally.","github_url":"https://github.com/unslothai/unsloth","owner":"unslothai","repo":"unsloth","owner_avatar_url":"https://avatars.githubusercontent.com/u/150920049?v=4","primary_language":"Python","stars":69621,"forks":6285,"topics":["agent","deepseek","fine-tuning","gemma","gemma3","gpt-oss","llama","llama3","llm","llms","mistral","openai","qwen","reinforcement-learning","self-hosted","text-to-speech","tts","ui","unsloth"],"archived":false,"github_pushed_at":"2026-08-06T06:01:56+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/unslothai-unsloth","markdown_url":"https://www.graphcanon.com/tools/unslothai-unsloth.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/unslothai-unsloth","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=unslothai-unsloth","shared_categories":["inference-serving"]}]}}