{"data":{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving","tool_count":406,"tools":[{"slug":"huggingface-transformers","name":"transformers","tagline":"Transformers: the model-definition framework for state-of-the-art machine learning models in text, vision, audio, and multimodal models","github_url":"https://github.com/huggingface/transformers","owner":"huggingface","repo":"transformers","owner_avatar_url":"https://avatars.githubusercontent.com/u/25720743?v=4","primary_language":"Python","stars":164121,"forks":34249,"topics":["audio","deep-learning","deepseek","gemma","glm","hacktoberfest","llm","machine-learning","model-hub","natural-language-processing","nlp","pretrained-models","python","pytorch","pytorch-transformers","qwen","speech-recognition","transformer","vlm"],"archived":false,"github_pushed_at":"2026-08-15T22:28:12+00:00","maintenance_label":"Very active","stars_delta_30d":1457,"url":"https://www.graphcanon.com/tools/huggingface-transformers","markdown_url":"https://www.graphcanon.com/tools/huggingface-transformers.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/huggingface-transformers","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=huggingface-transformers"},{"slug":"open-webui-open-webui","name":"open-webui","tagline":"User-friendly AI Interface (Supports Ollama, OpenAI API, ...)","github_url":"https://github.com/open-webui/open-webui","owner":"open-webui","repo":"open-webui","owner_avatar_url":"https://avatars.githubusercontent.com/u/158137808?v=4","primary_language":"Python","stars":148875,"forks":21676,"topics":["ai","llm","llm-ui","llm-webui","llms","mcp","ollama","ollama-webui","open-webui","openai","openapi","rag","self-hosted","ui","webui"],"archived":false,"github_pushed_at":"2026-08-15T07:10:16+00:00","maintenance_label":"Very active","stars_delta_30d":3224,"url":"https://www.graphcanon.com/tools/open-webui-open-webui","markdown_url":"https://www.graphcanon.com/tools/open-webui-open-webui.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/open-webui-open-webui","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=open-webui-open-webui"},{"slug":"ollama-ollama","name":"ollama","tagline":"Get up and running with various large language models using Ollama.","github_url":"https://github.com/ollama/ollama","owner":"ollama","repo":"ollama","owner_avatar_url":"https://avatars.githubusercontent.com/u/151674099?v=4","primary_language":"Go","stars":177524,"forks":17229,"topics":["deepseek","gemma","gemma3","glm","go","golang","gpt-oss","llama","llama3","llm","llms","minimax","mistral","ollama","qwen"],"archived":false,"github_pushed_at":"2026-07-31T23:59:29+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/ollama-ollama","markdown_url":"https://www.graphcanon.com/tools/ollama-ollama.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ollama-ollama","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ollama-ollama"},{"slug":"vllm-project-vllm","name":"vllm","tagline":"A high-throughput and memory-efficient inference and serving engine for LLMs","github_url":"https://github.com/vllm-project/vllm","owner":"vllm-project","repo":"vllm","owner_avatar_url":"https://avatars.githubusercontent.com/u/136984999?v=4","primary_language":"Python","stars":87847,"forks":20135,"topics":["amd","blackwell","cuda","deepseek","deepseek-v3","gpt","gpt-oss","inference","kimi","llama","llm","llm-serving","model-serving","moe","openai","pytorch","qwen","qwen3","tpu","transformer"],"archived":false,"github_pushed_at":"2026-08-01T11:55:36+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/vllm-project-vllm","markdown_url":"https://www.graphcanon.com/tools/vllm-project-vllm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/vllm-project-vllm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=vllm-project-vllm"},{"slug":"thedotmack-claude-mem","name":"claude-mem","tagline":"Persistent Context Across Sessions for Every Agent","github_url":"https://github.com/thedotmack/claude-mem","owner":"thedotmack","repo":"claude-mem","owner_avatar_url":"https://avatars.githubusercontent.com/u/683968?v=4","primary_language":"JavaScript","stars":91018,"forks":7953,"topics":["ai","ai-agents","ai-memory","anthropic","artificial-intelligence","chromadb","claude","claude-agent-sdk","claude-agents","claude-code","claude-code-plugin","claude-skills","embeddings","long-term-memory","mem0","memory-engine","openmemory","rag","sqlite","supermemory"],"archived":false,"github_pushed_at":"2026-08-17T15:46:17+00:00","maintenance_label":"Very active","stars_delta_30d":3316,"url":"https://www.graphcanon.com/tools/thedotmack-claude-mem","markdown_url":"https://www.graphcanon.com/tools/thedotmack-claude-mem.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/thedotmack-claude-mem","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=thedotmack-claude-mem"},{"slug":"unslothai-unsloth","name":"unsloth","tagline":"A web UI for training and running open models locally.","github_url":"https://github.com/unslothai/unsloth","owner":"unslothai","repo":"unsloth","owner_avatar_url":"https://avatars.githubusercontent.com/u/150920049?v=4","primary_language":"Python","stars":69621,"forks":6285,"topics":["agent","deepseek","fine-tuning","gemma","gemma3","gpt-oss","llama","llama3","llm","llms","mistral","openai","qwen","reinforcement-learning","self-hosted","text-to-speech","tts","ui","unsloth"],"archived":false,"github_pushed_at":"2026-08-06T06:01:56+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/unslothai-unsloth","markdown_url":"https://www.graphcanon.com/tools/unslothai-unsloth.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/unslothai-unsloth","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=unslothai-unsloth"},{"slug":"mintplex-labs-anything-llm","name":"anything-llm","tagline":"Self-hosted agent experience with deployment scripts for multiple environments","github_url":"https://github.com/Mintplex-Labs/anything-llm","owner":"Mintplex-Labs","repo":"anything-llm","owner_avatar_url":"https://avatars.githubusercontent.com/u/134426827?v=4","primary_language":"JavaScript","stars":64716,"forks":7132,"topics":["agent-computer","agent-harness","agent-orchestration","agentic-ai","ai-agents","computer-use","hermes-agent","llm","local-ai","localai","multimodal","no-code","open-claw","rag","self-hosted-ai","vector-database"],"archived":false,"github_pushed_at":"2026-08-13T22:31:08+00:00","maintenance_label":"Very active","stars_delta_30d":1371,"url":"https://www.graphcanon.com/tools/mintplex-labs-anything-llm","markdown_url":"https://www.graphcanon.com/tools/mintplex-labs-anything-llm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/mintplex-labs-anything-llm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=mintplex-labs-anything-llm"},{"slug":"jeecgboot-jeecgboot","name":"JeecgBoot","tagline":"AI低代码平台，实现快速生成前后端系统及模块","github_url":"https://github.com/jeecgboot/JeecgBoot","owner":"jeecgboot","repo":"JeecgBoot","owner_avatar_url":"https://avatars.githubusercontent.com/u/86360035?v=4","primary_language":"Java","stars":47405,"forks":16153,"topics":["activiti","agent","ai","antd","claude-code","cli","codegenerator","codex","flowable","langchain4j","llm","low-code","mcp","mybatis-plus","rag","skills","spring-ai","springboot","springcloud","vue3"],"archived":false,"github_pushed_at":"2026-08-14T06:09:59+00:00","maintenance_label":"Very active","stars_delta_30d":301,"url":"https://www.graphcanon.com/tools/jeecgboot-jeecgboot","markdown_url":"https://www.graphcanon.com/tools/jeecgboot-jeecgboot.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/jeecgboot-jeecgboot","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=jeecgboot-jeecgboot"},{"slug":"berriai-litellm","name":"litellm","tagline":"Python SDK and Proxy Server for calling multiple LLM APIs","github_url":"https://github.com/BerriAI/litellm","owner":"BerriAI","repo":"litellm","owner_avatar_url":"https://avatars.githubusercontent.com/u/121462774?v=4","primary_language":"Python","stars":55221,"forks":10231,"topics":["ai-gateway","anthropic","azure-openai","bedrock","gateway","langchain","litellm","llm","llm-gateway","llmops","mcp-gateway","openai","openai-proxy","rust","rust-ai","vertex-ai"],"archived":false,"github_pushed_at":"2026-08-01T05:53:28+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/berriai-litellm","markdown_url":"https://www.graphcanon.com/tools/berriai-litellm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/berriai-litellm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=berriai-litellm"},{"slug":"diegosouzapw-omniroute","name":"OmniRoute","tagline":"Free AI gateway with multi-provider support and token savings","github_url":"https://github.com/diegosouzapw/OmniRoute","owner":"diegosouzapw","repo":"OmniRoute","owner_avatar_url":"https://avatars.githubusercontent.com/u/8016841?v=4","primary_language":"TypeScript","stars":51233,"forks":6970,"topics":["a2a","ai-agents","ai-gateway","anthropic","claude","claude-code","cline","codex","copilot","cursor","deepseek","free-ai","gemini","kimi","llm-gateway","mcp","openai","openai-proxy","qwen","token-saver"],"archived":false,"github_pushed_at":"2026-08-19T21:44:14+00:00","maintenance_label":"Very active","stars_delta_30d":29928,"url":"https://www.graphcanon.com/tools/diegosouzapw-omniroute","markdown_url":"https://www.graphcanon.com/tools/diegosouzapw-omniroute.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/diegosouzapw-omniroute","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=diegosouzapw-omniroute"},{"slug":"ray-project-ray","name":"ray","tagline":"Ray is an AI compute engine with a core distributed runtime and AI Libraries for accelerating ML workloads.","github_url":"https://github.com/ray-project/ray","owner":"ray-project","repo":"ray","owner_avatar_url":"https://avatars.githubusercontent.com/u/22125274?v=4","primary_language":"Python","stars":43526,"forks":7929,"topics":["data-science","deep-learning","deployment","distributed","hyperparameter-optimization","hyperparameter-search","large-language-models","llm","llm-inference","llm-serving","machine-learning","optimization","parallel","python","pytorch","ray","reinforcement-learning","rllib","serving","tensorflow"],"archived":false,"github_pushed_at":"2026-08-16T00:26:16+00:00","maintenance_label":"Very active","stars_delta_30d":270,"url":"https://www.graphcanon.com/tools/ray-project-ray","markdown_url":"https://www.graphcanon.com/tools/ray-project-ray.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ray-project-ray","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ray-project-ray"},{"slug":"danny-avila-librechat","name":"LibreChat","tagline":"Enhanced ChatGPT Clone with extensive features and integrations for self-hosting","github_url":"https://github.com/danny-avila/LibreChat","owner":"danny-avila","repo":"LibreChat","owner_avatar_url":"https://avatars.githubusercontent.com/u/110412045?v=4","primary_language":"TypeScript","stars":41282,"forks":8495,"topics":["ai","anthropic","artifacts","aws","azure","chatgpt","chatgpt-clone","claude","clone","deepseek","gemini","google","gpt-5","librechat","mcp","o1","openai","responses-api","vision","webui"],"archived":false,"github_pushed_at":"2026-07-25T21:44:54+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/danny-avila-librechat","markdown_url":"https://www.graphcanon.com/tools/danny-avila-librechat.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/danny-avila-librechat","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=danny-avila-librechat"},{"slug":"sgl-project-sglang","name":"sglang","tagline":"High-performance serving framework for large language and multimodal models","github_url":"https://github.com/sgl-project/sglang","owner":"sgl-project","repo":"sglang","owner_avatar_url":"https://avatars.githubusercontent.com/u/147780389?v=4","primary_language":"Python","stars":31454,"forks":7720,"topics":["attention","blackwell","cuda","deepseek","diffusion","glm","gpt-oss","inference","llama","llm","minimax","moe","qwen","qwen-image","reinforcement-learning","transformer","vlm","wan"],"archived":false,"github_pushed_at":"2026-08-07T06:00:20+00:00","maintenance_label":"Very active","stars_delta_30d":1409,"url":"https://www.graphcanon.com/tools/sgl-project-sglang","markdown_url":"https://www.graphcanon.com/tools/sgl-project-sglang.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/sgl-project-sglang","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=sgl-project-sglang"},{"slug":"ggml-org-whisper-cpp","name":"whisper.cpp","tagline":"Port of OpenAI's Whisper model in C/C++ for speech-to-text inference","github_url":"https://github.com/ggml-org/whisper.cpp","owner":"ggml-org","repo":"whisper.cpp","owner_avatar_url":"https://avatars.githubusercontent.com/u/134263123?v=4","primary_language":"C++","stars":52501,"forks":5971,"topics":["inference","openai","speech-recognition","speech-to-text","transformer","whisper"],"archived":false,"github_pushed_at":"2026-07-31T07:11:28+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/ggml-org-whisper-cpp","markdown_url":"https://www.graphcanon.com/tools/ggml-org-whisper-cpp.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ggml-org-whisper-cpp","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ggml-org-whisper-cpp"},{"slug":"mlflow-mlflow","name":"mlflow","tagline":"AI engineering platform for debugging, evaluating, monitoring, and optimizing AI applications","github_url":"https://github.com/mlflow/mlflow","owner":"mlflow","repo":"mlflow","owner_avatar_url":"https://avatars.githubusercontent.com/u/39938107?v=4","primary_language":"Python","stars":27591,"forks":6189,"topics":["agentops","agents","ai","ai-governance","apache-spark","evaluation","langchain","llm-evaluation","llmops","machine-learning","ml","mlflow","mlops","model-management","observability","open-source","openai","prompt-engineering"],"archived":false,"github_pushed_at":"2026-08-20T00:54:28+00:00","maintenance_label":"Very active","stars_delta_30d":476,"url":"https://www.graphcanon.com/tools/mlflow-mlflow","markdown_url":"https://www.graphcanon.com/tools/mlflow-mlflow.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/mlflow-mlflow","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=mlflow-mlflow"},{"slug":"lutzroeder-netron","name":"netron","tagline":"Visualizer for neural network, deep learning and machine learning models","github_url":"https://github.com/lutzroeder/netron","owner":"lutzroeder","repo":"netron","owner_avatar_url":"https://avatars.githubusercontent.com/u/438516?v=4","primary_language":"JavaScript","stars":33302,"forks":3175,"topics":["ai","coreml","deep-learning","deeplearning","keras","machine-learning","machinelearning","ml","neural-network","numpy","onnx","pytorch","safetensors","tensorflow","tensorflow-lite","visualizer"],"archived":false,"github_pushed_at":"2026-08-02T16:49:54+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/lutzroeder-netron","markdown_url":"https://www.graphcanon.com/tools/lutzroeder-netron.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/lutzroeder-netron","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=lutzroeder-netron"},{"slug":"ruvnet-ruflo","name":"ruflo","tagline":"The leading agent meta-harness for intelligent multi-player swarms and autonomous workflows","github_url":"https://github.com/ruvnet/ruflo","owner":"ruvnet","repo":"ruflo","owner_avatar_url":"https://avatars.githubusercontent.com/u/2934394?v=4","primary_language":"TypeScript","stars":68322,"forks":8204,"topics":["agentic-ai","agentic-framework","agentic-workflow","agents","ai-agents","ai-assistant","ai-coding","ai-skills","autonomous-agents","claude-code","codex","harness","mcp-server","multi-agent","multi-agent-systems","npm","skills","swarm","swarm-intelligence","typescript"],"archived":false,"github_pushed_at":"2026-08-19T06:22:54+00:00","maintenance_label":"Very active","stars_delta_30d":3095,"url":"https://www.graphcanon.com/tools/ruvnet-ruflo","markdown_url":"https://www.graphcanon.com/tools/ruvnet-ruflo.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ruvnet-ruflo","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ruvnet-ruflo"},{"slug":"decolua-9router","name":"9router","tagline":"Unlimited FREE AI coding with auto-fallback and token savings","github_url":"https://github.com/decolua/9router","owner":"decolua","repo":"9router","owner_avatar_url":"https://avatars.githubusercontent.com/u/8282593?v=4","primary_language":"JavaScript","stars":25841,"forks":4635,"topics":["ai-agents","ai-gateway","anthropic","chatgpt","claude","claude-code","cline","codex","copilot","cursor","deepseek","free-ai","gemini","gemini-cli","llm","llm-gateway","openai","openai-proxy","qwen","token-saver"],"archived":false,"github_pushed_at":"2026-08-14T10:08:34+00:00","maintenance_label":"Very active","stars_delta_30d":3072,"url":"https://www.graphcanon.com/tools/decolua-9router","markdown_url":"https://www.graphcanon.com/tools/decolua-9router.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/decolua-9router","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=decolua-9router"},{"slug":"tencent-ncnn","name":"ncnn","tagline":"High-performance neural network inference framework optimized for mobile platforms","github_url":"https://github.com/Tencent/ncnn","owner":"Tencent","repo":"ncnn","owner_avatar_url":"https://avatars.githubusercontent.com/u/18461506?v=4","primary_language":"C++","stars":23644,"forks":4475,"topics":["android","arm-neon","artificial-intelligence","caffe","darknet","deep-learning","high-preformance","inference","ios","keras","mlir","mxnet","ncnn","neural-network","onnx","pytorch","riscv","simd","tensorflow","vulkan"],"archived":false,"github_pushed_at":"2026-08-04T09:18:11+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/tencent-ncnn","markdown_url":"https://www.graphcanon.com/tools/tencent-ncnn.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/tencent-ncnn","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=tencent-ncnn"},{"slug":"screenpipe-screenpipe","name":"screenpipe","tagline":"AI that records and analyzes everything you do, say, hear locally","github_url":"https://github.com/screenpipe/screenpipe","owner":"screenpipe","repo":"screenpipe","owner_avatar_url":"https://avatars.githubusercontent.com/u/259178917?v=4","primary_language":"Rust","stars":20534,"forks":2029,"topics":["agents","agi","ai","ai-memory","audio-recording","computer-vision","hermes","hermes-agent","llm","local-ai","local-first","machine-learning","mcp","multimodal","openclaw","privacy","rewind","screen-recording","speech-to-text","ycombinator"],"archived":false,"github_pushed_at":"2026-07-26T03:40:56+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/screenpipe-screenpipe","markdown_url":"https://www.graphcanon.com/tools/screenpipe-screenpipe.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/screenpipe-screenpipe","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=screenpipe-screenpipe"},{"slug":"liyupi-ai-guide","name":"ai-guide","tagline":"免费开放的AI知识共享平台","github_url":"https://github.com/liyupi/ai-guide","owner":"liyupi","repo":"ai-guide","owner_avatar_url":"https://avatars.githubusercontent.com/u/26037703?v=4","primary_language":"JavaScript","stars":18766,"forks":2127,"topics":["ai","artificial-intelligence","chatgpt","claude","codex","cursor","deep-learning","deepseek","gemini","generative-ai","gpt","llm","mcp","openai","python","rag","vibe-coding","vibecoding","vue","vuepress"],"archived":false,"github_pushed_at":"2026-08-07T14:50:37+00:00","maintenance_label":"Active","stars_delta_30d":1371,"url":"https://www.graphcanon.com/tools/liyupi-ai-guide","markdown_url":"https://www.graphcanon.com/tools/liyupi-ai-guide.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/liyupi-ai-guide","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=liyupi-ai-guide"},{"slug":"janhq-jan","name":"jan","tagline":"open source alternative to ChatGPT that runs offline locally","github_url":"https://github.com/janhq/jan","owner":"janhq","repo":"jan","owner_avatar_url":"https://avatars.githubusercontent.com/u/102363196?v=4","primary_language":"TypeScript","stars":44020,"forks":2971,"topics":["chatgpt","gpt","llamacpp","llm","localai","open-source","self-hosted","tauri"],"archived":false,"github_pushed_at":"2026-08-14T12:24:52+00:00","maintenance_label":"Very active","stars_delta_30d":426,"url":"https://www.graphcanon.com/tools/janhq-jan","markdown_url":"https://www.graphcanon.com/tools/janhq-jan.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/janhq-jan","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=janhq-jan"},{"slug":"stas00-ml-engineering","name":"ml-engineering","tagline":"Machine Learning Engineering Open Book","github_url":"https://github.com/stas00/ml-engineering","owner":"stas00","repo":"ml-engineering","owner_avatar_url":"https://avatars.githubusercontent.com/u/10676103?v=4","primary_language":"Python","stars":18632,"forks":1200,"topics":["ai","debugging","gpus","inference","large-language-models","llm","machine-learning","machine-learning-engineering","mlops","network","pytorch","scalability","slurm","storage","training","transformers"],"archived":false,"github_pushed_at":"2026-08-14T19:59:44+00:00","maintenance_label":"Very active","stars_delta_30d":216,"url":"https://www.graphcanon.com/tools/stas00-ml-engineering","markdown_url":"https://www.graphcanon.com/tools/stas00-ml-engineering.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/stas00-ml-engineering","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=stas00-ml-engineering"},{"slug":"lyogavin-airllm","name":"airllm","tagline":"AirLLM 70B inference with single 4GB GPU","github_url":"https://github.com/lyogavin/airllm","owner":"lyogavin","repo":"airllm","owner_avatar_url":"https://avatars.githubusercontent.com/u/1113905?v=4","primary_language":"Jupyter Notebook","stars":24183,"forks":2722,"topics":["chinese-llm","chinese-nlp","finetune","generative-ai","instruct-gpt","instruction-set","llama","llm","lora","open-models","open-source","open-source-models","qlora"],"archived":false,"github_pushed_at":"2026-07-23T08:29:43+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/lyogavin-airllm","markdown_url":"https://www.graphcanon.com/tools/lyogavin-airllm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/lyogavin-airllm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=lyogavin-airllm"},{"slug":"googlecloudplatform-generative-ai","name":"generative-ai","tagline":"Sample code and notebooks for Generative AI on Google Cloud, with Gemini Enterprise Agent Platform","github_url":"https://github.com/GoogleCloudPlatform/generative-ai","owner":"GoogleCloudPlatform","repo":"generative-ai","owner_avatar_url":"https://avatars.githubusercontent.com/u/2810941?v=4","primary_language":"Jupyter Notebook","stars":17594,"forks":4412,"topics":["agents","gcp","gemini","gemini-api","gen-ai","generative-ai","google","google-cloud","google-gemini","langchain","large-language-models","llm","vertex-ai","vertex-ai-gemini-api","vertexai"],"archived":false,"github_pushed_at":"2026-08-15T14:38:29+00:00","maintenance_label":"Very active","stars_delta_30d":247,"url":"https://www.graphcanon.com/tools/googlecloudplatform-generative-ai","markdown_url":"https://www.graphcanon.com/tools/googlecloudplatform-generative-ai.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/googlecloudplatform-generative-ai","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=googlecloudplatform-generative-ai"},{"slug":"openvinotoolkit-openvino","name":"openvino","tagline":"OpenVINO is an open source toolkit for optimizing and deploying AI inference.","github_url":"https://github.com/openvinotoolkit/openvino","owner":"openvinotoolkit","repo":"openvino","owner_avatar_url":"https://avatars.githubusercontent.com/u/55443902?v=4","primary_language":"C++","stars":10566,"forks":3286,"topics":["ai","computer-vision","deep-learning","deploy-ai","diffusion-models","generative-ai","good-first-issue","inference","llm-inference","natural-language-processing","nlp","openvino","optimize-ai","performance-boost","recommendation-system","speech-recognition","stable-diffusion","transformers","yolo"],"archived":false,"github_pushed_at":"2026-07-24T19:20:01+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/openvinotoolkit-openvino","markdown_url":"https://www.graphcanon.com/tools/openvinotoolkit-openvino.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/openvinotoolkit-openvino","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=openvinotoolkit-openvino"},{"slug":"deepspeedai-deepspeed","name":"DeepSpeed","tagline":"Deep learning optimization library for efficient distributed training and inference","github_url":"https://github.com/deepspeedai/DeepSpeed","owner":"deepspeedai","repo":"DeepSpeed","owner_avatar_url":"https://avatars.githubusercontent.com/u/74068820?v=4","primary_language":"Python","stars":42870,"forks":4920,"topics":["billion-parameters","compression","data-parallelism","deep-learning","gpu","inference","machine-learning","mixture-of-experts","model-parallelism","pipeline-parallelism","pytorch","trillion-parameters","zero"],"archived":false,"github_pushed_at":"2026-08-06T16:21:12+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/deepspeedai-deepspeed","markdown_url":"https://www.graphcanon.com/tools/deepspeedai-deepspeed.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/deepspeedai-deepspeed","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=deepspeedai-deepspeed"},{"slug":"jundot-omlx","name":"omlx","tagline":"LLM inference server with continuous batching and SSD caching for Apple Silicon","github_url":"https://github.com/jundot/omlx","owner":"jundot","repo":"omlx","owner_avatar_url":"https://avatars.githubusercontent.com/u/64250138?v=4","primary_language":"Python","stars":18679,"forks":1617,"topics":["apple-silicon","inference-server","llm","macos","mlx","openai-api"],"archived":false,"github_pushed_at":"2026-08-14T09:12:48+00:00","maintenance_label":"Very active","stars_delta_30d":839,"url":"https://www.graphcanon.com/tools/jundot-omlx","markdown_url":"https://www.graphcanon.com/tools/jundot-omlx.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/jundot-omlx","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=jundot-omlx"},{"slug":"pytorch-pytorch","name":"pytorch","tagline":"Tensors and Dynamic neural networks in Python with strong GPU acceleration","github_url":"https://github.com/pytorch/pytorch","owner":"pytorch","repo":"pytorch","owner_avatar_url":"https://avatars.githubusercontent.com/u/21003710?v=4","primary_language":"Python","stars":102144,"forks":28650,"topics":["autograd","deep-learning","gpu","machine-learning","neural-network","numpy","python","tensor"],"archived":false,"github_pushed_at":"2026-08-03T06:00:50+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/pytorch-pytorch","markdown_url":"https://www.graphcanon.com/tools/pytorch-pytorch.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/pytorch-pytorch","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=pytorch-pytorch"},{"slug":"bentoml-openllm","name":"OpenLLM","tagline":"Run any open-source LLMs as OpenAI compatible API endpoint in the cloud.","github_url":"https://github.com/bentoml/OpenLLM","owner":"bentoml","repo":"OpenLLM","owner_avatar_url":"https://avatars.githubusercontent.com/u/49176046?v=4","primary_language":"Python","stars":12454,"forks":828,"topics":["bentoml","fine-tuning","llama","llama2","llama3-1","llama3-2","llama3-2-vision","llm","llm-inference","llm-ops","llm-serving","llmops","mistral","mlops","model-inference","open-source-llm","openllm","vicuna"],"archived":false,"github_pushed_at":"2026-08-03T16:59:03+00:00","maintenance_label":"Very active","stars_delta_30d":66,"url":"https://www.graphcanon.com/tools/bentoml-openllm","markdown_url":"https://www.graphcanon.com/tools/bentoml-openllm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/bentoml-openllm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=bentoml-openllm"},{"slug":"xorbitsai-inference","name":"inference","tagline":"Unified production-ready inference API for various models","github_url":"https://github.com/xorbitsai/inference","owner":"xorbitsai","repo":"inference","owner_avatar_url":"https://avatars.githubusercontent.com/u/109655068?v=4","primary_language":"Python","stars":9470,"forks":851,"topics":["artificial-intelligence","chatglm","deployment","flan-t5","gemma","ggml","glm4","inference","llama","llama3","llamacpp","llm","machine-learning","mistral","openai-api","pytorch","qwen","vllm","whisper","wizardlm"],"archived":false,"github_pushed_at":"2026-08-02T05:44:21+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/xorbitsai-inference","markdown_url":"https://www.graphcanon.com/tools/xorbitsai-inference.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/xorbitsai-inference","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=xorbitsai-inference"},{"slug":"wangrongsheng-awesome-llm-resources","name":"awesome-LLM-resources","tagline":"Summary of the world's best LLM resources.","github_url":"https://github.com/WangRongsheng/awesome-LLM-resources","owner":"WangRongsheng","repo":"awesome-LLM-resources","owner_avatar_url":"https://avatars.githubusercontent.com/u/55651568?v=4","primary_language":null,"stars":8845,"forks":950,"topics":["awesome-list","book","course","large-language-models","llama","llm","mistral","openai","qwen","rag","retrieval-augmented-generation","webui"],"archived":false,"github_pushed_at":"2026-08-14T15:54:28+00:00","maintenance_label":"Very active","stars_delta_30d":142,"url":"https://www.graphcanon.com/tools/wangrongsheng-awesome-llm-resources","markdown_url":"https://www.graphcanon.com/tools/wangrongsheng-awesome-llm-resources.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/wangrongsheng-awesome-llm-resources","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=wangrongsheng-awesome-llm-resources"},{"slug":"oumi-ai-oumi","name":"oumi","tagline":"Easily fine-tune, evaluate and deploy open source LLMs/VLMs","github_url":"https://github.com/oumi-ai/oumi","owner":"oumi-ai","repo":"oumi","owner_avatar_url":"https://avatars.githubusercontent.com/u/167452922?v=4","primary_language":"Python","stars":9359,"forks":780,"topics":["dpo","evaluation","fine-tuning","gpt-oss","gpt-oss-120b","gpt-oss-20b","inference","llama","llms","sft","slms","vlms"],"archived":false,"github_pushed_at":"2026-07-24T05:44:23+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/oumi-ai-oumi","markdown_url":"https://www.graphcanon.com/tools/oumi-ai-oumi.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/oumi-ai-oumi","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=oumi-ai-oumi"},{"slug":"ai-dynamo-dynamo","name":"dynamo","tagline":"A Datacenter Scale Distributed Inference Serving Framework","github_url":"https://github.com/ai-dynamo/dynamo","owner":"ai-dynamo","repo":"dynamo","owner_avatar_url":"https://avatars.githubusercontent.com/u/201626793?v=4","primary_language":"Rust","stars":7575,"forks":1368,"topics":["diffusion","disaggregated-serving","kubernetes","llm-inference","omni","routing-engine","rust","sglang","tensorrt-llm","vllm"],"archived":false,"github_pushed_at":"2026-07-25T05:43:22+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/ai-dynamo-dynamo","markdown_url":"https://www.graphcanon.com/tools/ai-dynamo-dynamo.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ai-dynamo-dynamo","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ai-dynamo-dynamo"},{"slug":"runanywhereai-runanywhere-sdks","name":"runanywhere-sdks","tagline":"Production ready toolkit to run AI locally","github_url":"https://github.com/RunanywhereAI/runanywhere-sdks","owner":"RunanywhereAI","repo":"runanywhere-sdks","owner_avatar_url":"https://avatars.githubusercontent.com/u/220821781?v=4","primary_language":"C++","stars":10300,"forks":368,"topics":["android","apple-intelligence","cpp","diffusion-models","edge","flutter","inference","ios","kotlin","llamacpp","llm","multimodal","ollama","on-device-ai","react-native","swift","vlm","voice-ai","web","websdk"],"archived":false,"github_pushed_at":"2026-08-13T17:37:41+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/runanywhereai-runanywhere-sdks","markdown_url":"https://www.graphcanon.com/tools/runanywhereai-runanywhere-sdks.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/runanywhereai-runanywhere-sdks","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=runanywhereai-runanywhere-sdks"},{"slug":"funaudiollm-cosyvoice","name":"CosyVoice","tagline":"Multi-lingual large voice generation model with full-stack abilities for inference, training and deployment.","github_url":"https://github.com/FunAudioLLM/CosyVoice","owner":"FunAudioLLM","repo":"CosyVoice","owner_avatar_url":"https://avatars.githubusercontent.com/u/167062371?v=4","primary_language":"Python","stars":22373,"forks":2584,"topics":["audio-generation","cantonese","chatbot","chatgpt","chinese","cosyvoice","cross-lingual","english","fine-grained","fine-tuning","gpt-4o","japanese","korean","multi-lingual","natural-language-generation","python","text-to-speech","tts","voice-cloning"],"archived":false,"github_pushed_at":"2026-05-25T18:15:40+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/funaudiollm-cosyvoice","markdown_url":"https://www.graphcanon.com/tools/funaudiollm-cosyvoice.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/funaudiollm-cosyvoice","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=funaudiollm-cosyvoice"},{"slug":"fareedkhan-dev-train-llm-from-scratch","name":"train-llm-from-scratch","tagline":"A straightforward method for training your LLM from raw text to aligned model generation","github_url":"https://github.com/FareedKhan-dev/train-llm-from-scratch","owner":"FareedKhan-dev","repo":"train-llm-from-scratch","owner_avatar_url":"https://avatars.githubusercontent.com/u/63067900?v=4","primary_language":"Python","stars":9141,"forks":1264,"topics":["gemini","large-language-models","llm","openai","training","transformers"],"archived":false,"github_pushed_at":"2026-08-17T05:07:26+00:00","maintenance_label":"Very active","stars_delta_30d":765,"url":"https://www.graphcanon.com/tools/fareedkhan-dev-train-llm-from-scratch","markdown_url":"https://www.graphcanon.com/tools/fareedkhan-dev-train-llm-from-scratch.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/fareedkhan-dev-train-llm-from-scratch","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=fareedkhan-dev-train-llm-from-scratch"},{"slug":"mnfst-awesome-free-llm-apis","name":"awesome-free-llm-apis","tagline":"List of Permanent Free LLM API","github_url":"https://github.com/mnfst/awesome-free-llm-apis","owner":"mnfst","repo":"awesome-free-llm-apis","owner_avatar_url":"https://avatars.githubusercontent.com/u/8403534?v=4","primary_language":"JavaScript","stars":6532,"forks":630,"topics":["ai-agents","anthropic","awesome","awesome-list","gemini","llm","llm-router","llm-routing","ollama","openai","openclaw","openclaw-plugin","router"],"archived":false,"github_pushed_at":"2026-07-30T15:05:10+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/mnfst-awesome-free-llm-apis","markdown_url":"https://www.graphcanon.com/tools/mnfst-awesome-free-llm-apis.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/mnfst-awesome-free-llm-apis","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=mnfst-awesome-free-llm-apis"},{"slug":"katanemo-plano","name":"plano","tagline":"An AI-native proxy and data plane for agentic apps","github_url":"https://github.com/katanemo/plano","owner":"katanemo","repo":"plano","owner_avatar_url":"https://avatars.githubusercontent.com/u/112724757?v=4","primary_language":"Rust","stars":7004,"forks":482,"topics":["ai-gateway","ai-gateway-support","envoy","envoyproxy","gateway","generative-ai","llm-gateway","llm-inference","llm-proxy","llm-routing","llmops","llms","openai","prompt","proxy","proxy-server","routing"],"archived":false,"github_pushed_at":"2026-08-19T19:29:08+00:00","maintenance_label":"Very active","stars_delta_30d":127,"url":"https://www.graphcanon.com/tools/katanemo-plano","markdown_url":"https://www.graphcanon.com/tools/katanemo-plano.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/katanemo-plano","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=katanemo-plano"},{"slug":"bentoml-bentoml","name":"BentoML","tagline":"The easiest way to serve AI apps and models","github_url":"https://github.com/bentoml/BentoML","owner":"bentoml","repo":"BentoML","owner_avatar_url":"https://avatars.githubusercontent.com/u/49176046?v=4","primary_language":"Python","stars":8793,"forks":1010,"topics":["ai-inference","deep-learning","generative-ai","inference-platform","llm","llm-inference","llm-serving","llmops","machine-learning","ml-engineering","mlops","model-inference-service","model-serving","multimodal","python"],"archived":false,"github_pushed_at":"2026-08-03T17:00:21+00:00","maintenance_label":"Active","stars_delta_30d":65,"url":"https://www.graphcanon.com/tools/bentoml-bentoml","markdown_url":"https://www.graphcanon.com/tools/bentoml-bentoml.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/bentoml-bentoml","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=bentoml-bentoml"},{"slug":"tensorflow-serving","name":"serving","tagline":"A flexible, high-performance serving system for machine learning models","github_url":"https://github.com/tensorflow/serving","owner":"tensorflow","repo":"serving","owner_avatar_url":"https://avatars.githubusercontent.com/u/15658638?v=4","primary_language":"C++","stars":6359,"forks":2204,"topics":["cpp","deep-learning","deep-neural-networks","machine-learning","ml","neural-network","python","serving","tensorflow"],"archived":false,"github_pushed_at":"2026-07-30T07:02:43+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/tensorflow-serving","markdown_url":"https://www.graphcanon.com/tools/tensorflow-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/tensorflow-serving","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=tensorflow-serving"},{"slug":"ggml-org-llama-cpp","name":"llama.cpp","tagline":"LLM inference in C/C++","github_url":"https://github.com/ggml-org/llama.cpp","owner":"ggml-org","repo":"llama.cpp","owner_avatar_url":"https://avatars.githubusercontent.com/u/134263123?v=4","primary_language":"C++","stars":122941,"forks":21406,"topics":["ggml"],"archived":false,"github_pushed_at":"2026-08-07T05:28:54+00:00","maintenance_label":"Very active","stars_delta_30d":3353,"url":"https://www.graphcanon.com/tools/ggml-org-llama-cpp","markdown_url":"https://www.graphcanon.com/tools/ggml-org-llama-cpp.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ggml-org-llama-cpp","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ggml-org-llama-cpp"},{"slug":"alibaba-mnn","name":"MNN","tagline":"Blazing-fast, lightweight inference engine for high-performance on-device LLMs and Edge AI","github_url":"https://github.com/alibaba/MNN","owner":"alibaba","repo":"MNN","owner_avatar_url":"https://avatars.githubusercontent.com/u/1961952?v=4","primary_language":"C++","stars":15830,"forks":2398,"topics":["arm","convolution","deep-learning","embedded-devices","llm","machine-learning","ml","mnn","transformer","vulkan","winograd-algorithm"],"archived":false,"github_pushed_at":"2026-08-07T15:38:28+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/alibaba-mnn","markdown_url":"https://www.graphcanon.com/tools/alibaba-mnn.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/alibaba-mnn","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=alibaba-mnn"},{"slug":"internlm-lmdeploy","name":"lmdeploy","tagline":"Toolkit for compressing, deploying, and serving LLMs","github_url":"https://github.com/InternLM/lmdeploy","owner":"InternLM","repo":"lmdeploy","owner_avatar_url":"https://avatars.githubusercontent.com/u/135356492?v=4","primary_language":"Python","stars":7995,"forks":723,"topics":["codellama","cuda-kernels","deepspeed","fastertransformer","internlm","llama","llama2","llama3","llm","llm-inference","turbomind"],"archived":false,"github_pushed_at":"2026-08-06T09:17:57+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/internlm-lmdeploy","markdown_url":"https://www.graphcanon.com/tools/internlm-lmdeploy.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/internlm-lmdeploy","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=internlm-lmdeploy"},{"slug":"netflix-metaflow","name":"metaflow","tagline":"Build, Manage and Deploy AI/ML Systems","github_url":"https://github.com/Netflix/metaflow","owner":"Netflix","repo":"metaflow","owner_avatar_url":"https://avatars.githubusercontent.com/u/913567?v=4","primary_language":"Python","stars":10228,"forks":1330,"topics":["agents","ai","aws","azure","cost-optimization","datascience","distributed-training","gcp","generative-ai","high-performance-computing","kubernetes","llm","llmops","machine-learning","ml","ml-infrastructure","ml-platform","mlops","model-management","python"],"archived":false,"github_pushed_at":"2026-08-18T09:41:43+00:00","maintenance_label":"Very active","stars_delta_30d":38,"url":"https://www.graphcanon.com/tools/netflix-metaflow","markdown_url":"https://www.graphcanon.com/tools/netflix-metaflow.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/netflix-metaflow","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=netflix-metaflow"},{"slug":"mozilla-ai-llamafile","name":"llamafile","tagline":"Distribute and run LLMs with a single file.","github_url":"https://github.com/mozilla-ai/llamafile","owner":"mozilla-ai","repo":"llamafile","owner_avatar_url":"https://avatars.githubusercontent.com/u/129804596?v=4","primary_language":"C++","stars":25470,"forks":1530,"topics":["cross-platform","gguf","llama-cpp","local-ai","local-inference","local-llm","open-source-ai","single-file-executable","speech-to-text"],"archived":false,"github_pushed_at":"2026-07-27T14:21:26+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/mozilla-ai-llamafile","markdown_url":"https://www.graphcanon.com/tools/mozilla-ai-llamafile.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/mozilla-ai-llamafile","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=mozilla-ai-llamafile"},{"slug":"datawhalechina-llm-universe","name":"llm-universe","tagline":"面向小白开发者的LLM应用开发教程","github_url":"https://github.com/datawhalechina/llm-universe","owner":"datawhalechina","repo":"llm-universe","owner_avatar_url":"https://avatars.githubusercontent.com/u/46047812?v=4","primary_language":"Jupyter Notebook","stars":13803,"forks":1400,"topics":["langchain","rag"],"archived":false,"github_pushed_at":"2026-07-28T13:47:59+00:00","maintenance_label":"Active","stars_delta_30d":289,"url":"https://www.graphcanon.com/tools/datawhalechina-llm-universe","markdown_url":"https://www.graphcanon.com/tools/datawhalechina-llm-universe.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/datawhalechina-llm-universe","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=datawhalechina-llm-universe"},{"slug":"ther1d-shell-gpt","name":"shell_gpt","tagline":"A command-line productivity tool powered by AI large language models","github_url":"https://github.com/TheR1D/shell_gpt","owner":"TheR1D","repo":"shell_gpt","owner_avatar_url":"https://avatars.githubusercontent.com/u/16740832?v=4","primary_language":"Python","stars":12235,"forks":977,"topics":["chatgpt","cheat-sheet","cli","commands","gpt-3","gpt-4","gpt-5","linux","llama","llm","ollama","openai","productivity","python","shell","terminal"],"archived":false,"github_pushed_at":"2026-07-02T06:03:01+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/ther1d-shell-gpt","markdown_url":"https://www.graphcanon.com/tools/ther1d-shell-gpt.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ther1d-shell-gpt","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ther1d-shell-gpt"},{"slug":"liguodongiot-llm-action","name":"llm-action","tagline":"Aims to share large model technology principles and practical experience (large model engineering, application implementation)","github_url":"https://github.com/liguodongiot/llm-action","owner":"liguodongiot","repo":"llm-action","owner_avatar_url":"https://avatars.githubusercontent.com/u/13220186?v=4","primary_language":"HTML","stars":24898,"forks":2842,"topics":["llm","llm-inference","llm-serving","llm-training","llmops"],"archived":false,"github_pushed_at":"2026-07-19T13:13:31+00:00","maintenance_label":"Active","stars_delta_30d":162,"url":"https://www.graphcanon.com/tools/liguodongiot-llm-action","markdown_url":"https://www.graphcanon.com/tools/liguodongiot-llm-action.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/liguodongiot-llm-action","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=liguodongiot-llm-action"},{"slug":"maximhq-bifrost","name":"bifrost","tagline":"Fast Enterprise AI Gateway with Adaptive Load Balancer and Guardrails","github_url":"https://github.com/maximhq/bifrost","owner":"maximhq","repo":"bifrost","owner_avatar_url":"https://avatars.githubusercontent.com/u/139708451?v=4","primary_language":"Go","stars":7449,"forks":1073,"topics":["ai-gateway","gateway","gateway-services","generative-ai","guardrails","llm","llm-cost","llm-gateway","llm-observability","llmops","load-balancing","mcp-client","mcp-gateway","mcp-server","model-router","token-management"],"archived":false,"github_pushed_at":"2026-08-20T11:57:33+00:00","maintenance_label":"Very active","stars_delta_30d":812,"url":"https://www.graphcanon.com/tools/maximhq-bifrost","markdown_url":"https://www.graphcanon.com/tools/maximhq-bifrost.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/maximhq-bifrost","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=maximhq-bifrost"},{"slug":"nvidia-tensorrt-llm","name":"TensorRT-LLM","tagline":"Python API for defining and optimizing Large Language Models (LLMs) on NVIDIA GPUs","github_url":"https://github.com/NVIDIA/TensorRT-LLM","owner":"NVIDIA","repo":"TensorRT-LLM","owner_avatar_url":"https://avatars.githubusercontent.com/u/1728152?v=4","primary_language":"Python","stars":14317,"forks":2641,"topics":["blackwell","cuda","llm-serving","moe","pytorch"],"archived":false,"github_pushed_at":"2026-08-07T05:40:26+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/nvidia-tensorrt-llm","markdown_url":"https://www.graphcanon.com/tools/nvidia-tensorrt-llm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/nvidia-tensorrt-llm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=nvidia-tensorrt-llm"},{"slug":"awslabs-mcp","name":"mcp","tagline":"Open source MCP Servers for AWS","github_url":"https://github.com/awslabs/mcp","owner":"awslabs","repo":"mcp","owner_avatar_url":"https://avatars.githubusercontent.com/u/3299148?v=4","primary_language":"Python","stars":9503,"forks":1661,"topics":["aws","mcp","mcp-client","mcp-clients","mcp-host","mcp-server","mcp-servers","mcp-tools","modelcontextprotocol"],"archived":false,"github_pushed_at":"2026-07-25T05:55:41+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/awslabs-mcp","markdown_url":"https://www.graphcanon.com/tools/awslabs-mcp.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/awslabs-mcp","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=awslabs-mcp"},{"slug":"lightning-ai-litgpt","name":"litgpt","tagline":"High-performance LLMs with recipes for pretraining, finetuning and deployment","github_url":"https://github.com/Lightning-AI/litgpt","owner":"Lightning-AI","repo":"litgpt","owner_avatar_url":"https://avatars.githubusercontent.com/u/58386951?v=4","primary_language":"Python","stars":13605,"forks":1483,"topics":["ai","artificial-intelligence","deep-learning","large-language-models","llm","llm-inference","llms"],"archived":false,"github_pushed_at":"2026-07-20T10:24:12+00:00","maintenance_label":"Active","stars_delta_30d":137,"url":"https://www.graphcanon.com/tools/lightning-ai-litgpt","markdown_url":"https://www.graphcanon.com/tools/lightning-ai-litgpt.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/lightning-ai-litgpt","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=lightning-ai-litgpt"},{"slug":"datawhalechina-self-llm","name":"self-llm","tagline":"A guide for fine-tuning and deploying open-source large language models tailored for a Chinese audience on Linux.","github_url":"https://github.com/datawhalechina/self-llm","owner":"datawhalechina","repo":"self-llm","owner_avatar_url":"https://avatars.githubusercontent.com/u/46047812?v=4","primary_language":"Jupyter Notebook","stars":31722,"forks":3082,"topics":["chatglm","chatglm3","gemma-2b-it","glm-4","internlm2","llama3","llm","lora","minicpm","q-wen","qwen","qwen1-5","qwen2"],"archived":false,"github_pushed_at":"2026-07-30T01:58:34+00:00","maintenance_label":"Active","stars_delta_30d":412,"url":"https://www.graphcanon.com/tools/datawhalechina-self-llm","markdown_url":"https://www.graphcanon.com/tools/datawhalechina-self-llm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/datawhalechina-self-llm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=datawhalechina-self-llm"},{"slug":"simonw-llm","name":"llm","tagline":"Access large language models from the command-line","github_url":"https://github.com/simonw/llm","owner":"simonw","repo":"llm","owner_avatar_url":"https://avatars.githubusercontent.com/u/9599?v=4","primary_language":"Python","stars":12324,"forks":939,"topics":["ai","llms","openai"],"archived":false,"github_pushed_at":"2026-08-05T14:29:09+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/simonw-llm","markdown_url":"https://www.graphcanon.com/tools/simonw-llm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/simonw-llm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=simonw-llm"},{"slug":"lightning-ai-pytorch-lightning","name":"pytorch-lightning","tagline":"Pretrain, finetune ANY AI model of ANY size on 1 or 10,000+ GPUs with zero code changes.","github_url":"https://github.com/Lightning-AI/pytorch-lightning","owner":"Lightning-AI","repo":"pytorch-lightning","owner_avatar_url":"https://avatars.githubusercontent.com/u/58386951?v=4","primary_language":"Python","stars":31267,"forks":3768,"topics":["ai","artificial-intelligence","data-science","deep-learning","machine-learning","python","pytorch"],"archived":false,"github_pushed_at":"2026-08-03T01:13:09+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/lightning-ai-pytorch-lightning","markdown_url":"https://www.graphcanon.com/tools/lightning-ai-pytorch-lightning.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/lightning-ai-pytorch-lightning","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=lightning-ai-pytorch-lightning"},{"slug":"bitsandbytes-foundation-bitsandbytes","name":"bitsandbytes","tagline":"Large language model quantization toolkit for PyTorch.","github_url":"https://github.com/bitsandbytes-foundation/bitsandbytes","owner":"bitsandbytes-foundation","repo":"bitsandbytes","owner_avatar_url":"https://avatars.githubusercontent.com/u/175231607?v=4","primary_language":"Python","stars":8385,"forks":900,"topics":["llm","machine-learning","pytorch","qlora","quantization"],"archived":false,"github_pushed_at":"2026-07-29T18:27:51+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/bitsandbytes-foundation-bitsandbytes","markdown_url":"https://www.graphcanon.com/tools/bitsandbytes-foundation-bitsandbytes.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/bitsandbytes-foundation-bitsandbytes","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=bitsandbytes-foundation-bitsandbytes"},{"slug":"arvinlovegood-go-stock","name":"go-stock","tagline":"AI-empowered stock analysis tool for multiple markets","github_url":"https://github.com/ArvinLovegood/go-stock","owner":"ArvinLovegood","repo":"go-stock","owner_avatar_url":"https://avatars.githubusercontent.com/u/7401917?v=4","primary_language":"Go","stars":7236,"forks":1268,"topics":["ai-tools","deepseek","golang","lmstudio","naiveui","ollama","openai","stock","wails"],"archived":false,"github_pushed_at":"2026-08-07T22:39:50+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/arvinlovegood-go-stock","markdown_url":"https://www.graphcanon.com/tools/arvinlovegood-go-stock.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/arvinlovegood-go-stock","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=arvinlovegood-go-stock"},{"slug":"zylon-ai-private-gpt","name":"private-gpt","tagline":"Complete API layer for private AI applications on local models","github_url":"https://github.com/zylon-ai/private-gpt","owner":"zylon-ai","repo":"private-gpt","owner_avatar_url":"https://avatars.githubusercontent.com/u/143802295?v=4","primary_language":"Python","stars":57415,"forks":7607,"topics":["ai","ai-tools","on-premise"],"archived":false,"github_pushed_at":"2026-08-06T13:41:08+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/zylon-ai-private-gpt","markdown_url":"https://www.graphcanon.com/tools/zylon-ai-private-gpt.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/zylon-ai-private-gpt","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=zylon-ai-private-gpt"},{"slug":"kalyanks-nlp-llm-engineer-toolkit","name":"llm-engineer-toolkit","tagline":"A curated list of over 120 LLM libraries categorized.","github_url":"https://github.com/KalyanKS-NLP/llm-engineer-toolkit","owner":"KalyanKS-NLP","repo":"llm-engineer-toolkit","owner_avatar_url":"https://avatars.githubusercontent.com/u/202506543?v=4","primary_language":null,"stars":10767,"forks":1682,"topics":["ai-engineer","generative-ai","large-language-models","llm-engineer","llms"],"archived":false,"github_pushed_at":"2026-08-16T13:05:43+00:00","maintenance_label":"Very active","stars_delta_30d":106,"url":"https://www.graphcanon.com/tools/kalyanks-nlp-llm-engineer-toolkit","markdown_url":"https://www.graphcanon.com/tools/kalyanks-nlp-llm-engineer-toolkit.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/kalyanks-nlp-llm-engineer-toolkit","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=kalyanks-nlp-llm-engineer-toolkit"}]}}