{"data":{"node":{"slug":"eastriverlee-llm-swift","name":"LLM.swift","tagline":"LLM.swift enables local interaction with large language models for multiple Apple platforms.","github_url":"https://github.com/eastriverlee/LLM.swift","owner":"eastriverlee","repo":"LLM.swift","owner_avatar_url":"https://avatars.githubusercontent.com/u/43613000?v=4","primary_language":"Swift","stars":871,"forks":124,"topics":["gguf","ios","llm","llm-inference","macos","swift","tvos","visionos","watchos"],"archived":false,"github_pushed_at":"2026-07-19T06:04:49+00:00","maintenance_label":"Steady","stars_delta_30d":6,"url":"https://www.graphcanon.com/tools/eastriverlee-llm-swift","markdown_url":"https://www.graphcanon.com/tools/eastriverlee-llm-swift.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/eastriverlee-llm-swift","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=eastriverlee-llm-swift"},"categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"}],"tags":[{"slug":"gguf","name":"gguf"},{"slug":"ios","name":"ios"},{"slug":"llm","name":"llm"},{"slug":"llm-inference","name":"llm-inference"},{"slug":"macos","name":"macos"},{"slug":"swift","name":"swift"},{"slug":"tvos","name":"tvos"},{"slug":"visionos","name":"visionos"}],"edges":[],"neighbours":[{"slug":"ggml-org-llama-cpp","name":"llama.cpp","tagline":"LLM inference in C/C++","github_url":"https://github.com/ggml-org/llama.cpp","owner":"ggml-org","repo":"llama.cpp","owner_avatar_url":"https://avatars.githubusercontent.com/u/134263123?v=4","primary_language":"C++","stars":122941,"forks":21406,"topics":["ggml"],"archived":false,"github_pushed_at":"2026-08-07T05:28:54+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/ggml-org-llama-cpp","markdown_url":"https://www.graphcanon.com/tools/ggml-org-llama-cpp.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ggml-org-llama-cpp","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ggml-org-llama-cpp","shared_categories":["inference-serving"]},{"slug":"vllm-project-vllm","name":"vllm","tagline":"A high-throughput and memory-efficient inference and serving engine for LLMs","github_url":"https://github.com/vllm-project/vllm","owner":"vllm-project","repo":"vllm","owner_avatar_url":"https://avatars.githubusercontent.com/u/136984999?v=4","primary_language":"Python","stars":87847,"forks":20135,"topics":["amd","blackwell","cuda","deepseek","deepseek-v3","gpt","gpt-oss","inference","kimi","llama","llm","llm-serving","model-serving","moe","openai","pytorch","qwen","qwen3","tpu","transformer"],"archived":false,"github_pushed_at":"2026-08-01T11:55:36+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/vllm-project-vllm","markdown_url":"https://www.graphcanon.com/tools/vllm-project-vllm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/vllm-project-vllm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=vllm-project-vllm","shared_categories":["inference-serving"]},{"slug":"nomic-ai-gpt4all","name":"gpt4all","tagline":"Run Local LLMs on Any Device","github_url":"https://github.com/nomic-ai/gpt4all","owner":"nomic-ai","repo":"gpt4all","owner_avatar_url":"https://avatars.githubusercontent.com/u/102670180?v=4","primary_language":"C++","stars":77393,"forks":8296,"topics":["ai-chat","llm-inference"],"archived":false,"github_pushed_at":"2025-05-27T20:05:19+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/nomic-ai-gpt4all","markdown_url":"https://www.graphcanon.com/tools/nomic-ai-gpt4all.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/nomic-ai-gpt4all","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=nomic-ai-gpt4all","shared_categories":["inference-serving"]},{"slug":"jundot-omlx","name":"omlx","tagline":"LLM inference server with continuous batching and SSD caching for Apple Silicon","github_url":"https://github.com/jundot/omlx","owner":"jundot","repo":"omlx","owner_avatar_url":"https://avatars.githubusercontent.com/u/64250138?v=4","primary_language":"Python","stars":18679,"forks":1617,"topics":["apple-silicon","inference-server","llm","macos","mlx","openai-api"],"archived":false,"github_pushed_at":"2026-08-14T09:12:48+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/jundot-omlx","markdown_url":"https://www.graphcanon.com/tools/jundot-omlx.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/jundot-omlx","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=jundot-omlx","shared_categories":["inference-serving"]},{"slug":"lightning-ai-litgpt","name":"litgpt","tagline":"High-performance LLMs with recipes for pretraining, finetuning and deployment","github_url":"https://github.com/Lightning-AI/litgpt","owner":"Lightning-AI","repo":"litgpt","owner_avatar_url":"https://avatars.githubusercontent.com/u/58386951?v=4","primary_language":"Python","stars":13605,"forks":1483,"topics":["ai","artificial-intelligence","deep-learning","large-language-models","llm","llm-inference","llms"],"archived":false,"github_pushed_at":"2026-07-20T10:24:12+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/lightning-ai-litgpt","markdown_url":"https://www.graphcanon.com/tools/lightning-ai-litgpt.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/lightning-ai-litgpt","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=lightning-ai-litgpt","shared_categories":["inference-serving"]},{"slug":"kalyanks-nlp-llm-engineer-toolkit","name":"llm-engineer-toolkit","tagline":"A curated list of over 120 LLM libraries categorized.","github_url":"https://github.com/KalyanKS-NLP/llm-engineer-toolkit","owner":"KalyanKS-NLP","repo":"llm-engineer-toolkit","owner_avatar_url":"https://avatars.githubusercontent.com/u/202506543?v=4","primary_language":null,"stars":10767,"forks":1682,"topics":["ai-engineer","generative-ai","large-language-models","llm-engineer","llms"],"archived":false,"github_pushed_at":"2026-08-16T13:05:43+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/kalyanks-nlp-llm-engineer-toolkit","markdown_url":"https://www.graphcanon.com/tools/kalyanks-nlp-llm-engineer-toolkit.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/kalyanks-nlp-llm-engineer-toolkit","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=kalyanks-nlp-llm-engineer-toolkit","shared_categories":["inference-serving"]},{"slug":"argmaxinc-argmax-oss-swift","name":"argmax-oss-swift","tagline":"On-device Speech AI for Apple Silicon","github_url":"https://github.com/argmaxinc/argmax-oss-swift","owner":"argmaxinc","repo":"argmax-oss-swift","owner_avatar_url":"https://avatars.githubusercontent.com/u/150409474?v=4","primary_language":"Swift","stars":6294,"forks":591,"topics":["inference","ios","macos","pyannote","qwen3-tts","speaker-diarization","speakerkit","speech-recognition","speech-to-text","swift","text-to-speech","transformers","ttskit","whisper","whisperkit"],"archived":false,"github_pushed_at":"2026-07-28T23:05:28+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/argmaxinc-argmax-oss-swift","markdown_url":"https://www.graphcanon.com/tools/argmaxinc-argmax-oss-swift.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/argmaxinc-argmax-oss-swift","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=argmaxinc-argmax-oss-swift","shared_categories":[]},{"slug":"raullenchai-rapid-mlx","name":"Rapid-MLX","tagline":"Fast local AI engine for Apple Silicon","github_url":"https://github.com/raullenchai/Rapid-MLX","owner":"raullenchai","repo":"Rapid-MLX","owner_avatar_url":"https://avatars.githubusercontent.com/u/989846?v=4","primary_language":"Python","stars":3391,"forks":388,"topics":["apple-silicon","claude-code","cursor","deepseek","fastapi","hacktoberfest","inference","llm","local-llm","m1","m2","m3","macos","mlx","ollama-alternative","openai-api","python","qwen","tool-calling"],"archived":false,"github_pushed_at":"2026-08-01T23:14:51+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/raullenchai-rapid-mlx","markdown_url":"https://www.graphcanon.com/tools/raullenchai-rapid-mlx.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/raullenchai-rapid-mlx","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=raullenchai-rapid-mlx","shared_categories":["inference-serving"]},{"slug":"kevinhermawan-ollamac","name":"Ollamac","tagline":"Mac app for Ollama","github_url":"https://github.com/kevinhermawan/Ollamac","owner":"kevinhermawan","repo":"Ollamac","owner_avatar_url":"https://avatars.githubusercontent.com/u/84965338?v=4","primary_language":"Swift","stars":1912,"forks":99,"topics":["ai","llama3","llm","llma","llma2","llms","macos","mistral","mixtral","ollama"],"archived":false,"github_pushed_at":"2025-03-12T22:28:22+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/kevinhermawan-ollamac","markdown_url":"https://www.graphcanon.com/tools/kevinhermawan-ollamac.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/kevinhermawan-ollamac","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=kevinhermawan-ollamac","shared_categories":[]},{"slug":"waybarrios-vllm-mlx","name":"vllm-mlx","tagline":"Server for LLMs and vision-language models compatible with Apple Silicon","github_url":"https://github.com/waybarrios/vllm-mlx","owner":"waybarrios","repo":"vllm-mlx","owner_avatar_url":"https://avatars.githubusercontent.com/u/6794828?v=4","primary_language":"Python","stars":1472,"forks":205,"topics":["anthropic","apple-silicon","audio-processing","claude-code","computer-vision","image-understanding","inference","llm","machine-learning","macos","mllm","mlx","multimodal-ai","speech-to-text","stt","text-to-speech","tts","video-understanding","vision-language-model","vllm"],"archived":false,"github_pushed_at":"2026-06-28T20:18:31+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/waybarrios-vllm-mlx","markdown_url":"https://www.graphcanon.com/tools/waybarrios-vllm-mlx.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/waybarrios-vllm-mlx","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=waybarrios-vllm-mlx","shared_categories":["inference-serving"]},{"slug":"rudrankriyam-foundation-models-framework-lab","name":"Foundation-Models-Framework-Lab","tagline":"A practical lab for building, testing, and evaluating apps with Apple's Foundation Models framework","github_url":"https://github.com/rudrankriyam/Foundation-Models-Framework-Lab","owner":"rudrankriyam","repo":"Foundation-Models-Framework-Lab","owner_avatar_url":"https://avatars.githubusercontent.com/u/30552772?v=4","primary_language":"Swift","stars":1163,"forks":69,"topics":["ai","apple-foundation-models","apple-intelligence","foundation-models","foundation-models-framework","generative-ai","healthkit","ios","large-language-models","llm","macos","multilingual","on-device-ai","rag","speech-recognition","swift","swiftui","text-to-speech","tool-calling","xcode"],"archived":false,"github_pushed_at":"2026-07-20T21:06:23+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/rudrankriyam-foundation-models-framework-lab","markdown_url":"https://www.graphcanon.com/tools/rudrankriyam-foundation-models-framework-lab.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/rudrankriyam-foundation-models-framework-lab","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=rudrankriyam-foundation-models-framework-lab","shared_categories":[]},{"slug":"stoyan-stoyanov-llmflows","name":"llmflows","tagline":"Simple Explicit Transparent LLM Apps","github_url":"https://github.com/stoyan-stoyanov/llmflows","owner":"stoyan-stoyanov","repo":"llmflows","owner_avatar_url":"https://avatars.githubusercontent.com/u/14061867?v=4","primary_language":"Python","stars":707,"forks":35,"topics":["ai","chatgpt","gpt-4","llm","llm-inference","llmops","llms","machine-learning","openai","prompt-engineering","python","question-answering","vector-database"],"archived":false,"github_pushed_at":"2025-02-20T16:53:45+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/stoyan-stoyanov-llmflows","markdown_url":"https://www.graphcanon.com/tools/stoyan-stoyanov-llmflows.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/stoyan-stoyanov-llmflows","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=stoyan-stoyanov-llmflows","shared_categories":["inference-serving"]}]}}