{"data":{"node":{"slug":"basetenlabs-truss","name":"truss","tagline":"The simplest way to serve AI/ML models in production","github_url":"https://github.com/basetenlabs/truss","owner":"basetenlabs","repo":"truss","owner_avatar_url":"https://avatars.githubusercontent.com/u/54861414?v=4","primary_language":"Python","stars":1203,"forks":126,"topics":["artificial-intelligence","easy-to-use","falcon","inference-api","inference-server","machine-learning","model-serving","open-source","packaging","stable-diffusion","whisper","wizardlm"],"archived":false,"github_pushed_at":"2026-09-18T22:12:03+00:00","maintenance_label":"Very active","stars_delta_30d":15,"url":"https://www.graphcanon.com/tools/basetenlabs-truss","markdown_url":"https://www.graphcanon.com/tools/basetenlabs-truss.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/basetenlabs-truss","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=basetenlabs-truss"},"categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"}],"tags":[{"slug":"artificial-intelligence","name":"artificial-intelligence"},{"slug":"easy-to-use","name":"easy-to-use"},{"slug":"falcon","name":"falcon"},{"slug":"inference-api","name":"inference-api"},{"slug":"inference-server","name":"inference-server"},{"slug":"machine-learning","name":"machine-learning"},{"slug":"model-serving","name":"model-serving"},{"slug":"open-source","name":"open-source"}],"edges":[],"neighbours":[{"slug":"bentoml-bentoml","name":"BentoML","tagline":"The easiest way to serve AI apps and models","github_url":"https://github.com/bentoml/BentoML","owner":"bentoml","repo":"BentoML","owner_avatar_url":"https://avatars.githubusercontent.com/u/49176046?v=4","primary_language":"Python","stars":8847,"forks":1032,"topics":["ai-inference","deep-learning","generative-ai","inference-platform","llm","llm-inference","llm-serving","llmops","machine-learning","ml-engineering","mlops","model-inference-service","model-serving","multimodal","python"],"archived":false,"github_pushed_at":"2026-09-07T17:43:01+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/bentoml-bentoml","markdown_url":"https://www.graphcanon.com/tools/bentoml-bentoml.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/bentoml-bentoml","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=bentoml-bentoml","shared_categories":["inference-serving"]},{"slug":"ai-dynamo-dynamo","name":"dynamo","tagline":"A Datacenter Scale Distributed Inference Serving Framework","github_url":"https://github.com/ai-dynamo/dynamo","owner":"ai-dynamo","repo":"dynamo","owner_avatar_url":"https://avatars.githubusercontent.com/u/201626793?v=4","primary_language":"Rust","stars":8121,"forks":1601,"topics":["diffusion","disaggregated-serving","kubernetes","llm-inference","omni","routing-engine","rust","sglang","tensorrt-llm","vllm"],"archived":false,"github_pushed_at":"2026-09-19T17:01:54+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/ai-dynamo-dynamo","markdown_url":"https://www.graphcanon.com/tools/ai-dynamo-dynamo.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ai-dynamo-dynamo","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ai-dynamo-dynamo","shared_categories":["inference-serving"]},{"slug":"ericlbuehler-mistral-rs","name":"mistral.rs","tagline":"Fast flexible LLM inference","github_url":"https://github.com/EricLBuehler/mistral.rs","owner":"EricLBuehler","repo":"mistral.rs","owner_avatar_url":"https://avatars.githubusercontent.com/u/65165915?v=4","primary_language":"Rust","stars":7660,"forks":690,"topics":["llm","rust","uqff"],"archived":false,"github_pushed_at":"2026-09-06T20:01:45+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/ericlbuehler-mistral-rs","markdown_url":"https://www.graphcanon.com/tools/ericlbuehler-mistral-rs.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ericlbuehler-mistral-rs","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ericlbuehler-mistral-rs","shared_categories":["inference-serving"]},{"slug":"tensorflow-serving","name":"serving","tagline":"Flexible, high-performance serving system for machine learning models","github_url":"https://github.com/tensorflow/serving","owner":"tensorflow","repo":"serving","owner_avatar_url":"https://avatars.githubusercontent.com/u/15658638?v=4","primary_language":"C++","stars":6362,"forks":2207,"topics":["cpp","deep-learning","deep-neural-networks","machine-learning","ml","neural-network","python","serving","tensorflow"],"archived":false,"github_pushed_at":"2026-09-16T22:46:53+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/tensorflow-serving","markdown_url":"https://www.graphcanon.com/tools/tensorflow-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/tensorflow-serving","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=tensorflow-serving","shared_categories":["inference-serving"]},{"slug":"kserve-kserve","name":"kserve","tagline":"Standardized Distributed Generative and Predictive AI Inference Platform for Scalable, Multi-Framework Deployment on Kubernetes","github_url":"https://github.com/kserve/kserve","owner":"kserve","repo":"kserve","owner_avatar_url":"https://avatars.githubusercontent.com/u/83512434?v=4","primary_language":"Go","stars":5970,"forks":1689,"topics":["artificial-intelligence","cncf","genai","hacktoberfest","istio","k8s","knative","kserve","kubeflow","kubernetes","llm-inference","machine-learning","mlops","model-interpretability","model-serving","pytorch","service-mesh","tensorflow","vllm","xgboost"],"archived":false,"github_pushed_at":"2026-09-19T13:08:33+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/kserve-kserve","markdown_url":"https://www.graphcanon.com/tools/kserve-kserve.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/kserve-kserve","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=kserve-kserve","shared_categories":["inference-serving"]},{"slug":"vllm-project-semantic-router","name":"semantic-router","tagline":"System level intelligent runtime for Mixture-of-Models across edge, data center and cloud","github_url":"https://github.com/vllm-project/semantic-router","owner":"vllm-project","repo":"semantic-router","owner_avatar_url":"https://avatars.githubusercontent.com/u/136984999?v=4","primary_language":"Go","stars":5869,"forks":957,"topics":["ai-gateway","guardrails","inference","kubernetes","llm","llmrouter","mixture-of-models","pytorch","semantic-router","transformer","vllm"],"archived":false,"github_pushed_at":"2026-09-19T18:18:46+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/vllm-project-semantic-router","markdown_url":"https://www.graphcanon.com/tools/vllm-project-semantic-router.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/vllm-project-semantic-router","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=vllm-project-semantic-router","shared_categories":["inference-serving"]},{"slug":"seldonio-seldon-core","name":"seldon-core","tagline":"An MLOps framework to package, deploy, monitor and manage thousands of production machine learning models","github_url":"https://github.com/SeldonIO/seldon-core","owner":"SeldonIO","repo":"seldon-core","owner_avatar_url":"https://avatars.githubusercontent.com/u/10297834?v=4","primary_language":"Go","stars":4780,"forks":868,"topics":["aiops","deployment","kubernetes","machine-learning","machine-learning-operations","mlops","production-machine-learning","serving"],"archived":false,"github_pushed_at":"2026-03-23T11:39:54+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/seldonio-seldon-core","markdown_url":"https://www.graphcanon.com/tools/seldonio-seldon-core.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/seldonio-seldon-core","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=seldonio-seldon-core","shared_categories":["inference-serving"]},{"slug":"pytorch-serve","name":"serve","tagline":"Serve, optimize and scale PyTorch models in production","github_url":"https://github.com/pytorch/serve","owner":"pytorch","repo":"serve","owner_avatar_url":"https://avatars.githubusercontent.com/u/21003710?v=4","primary_language":"Java","stars":4344,"forks":880,"topics":["cpu","deep-learning","docker","gpu","kubernetes","machine-learning","metrics","mlops","optimization","pytorch","serving"],"archived":true,"github_pushed_at":"2025-08-06T19:17:08+00:00","maintenance_label":"Archived","url":"https://www.graphcanon.com/tools/pytorch-serve","markdown_url":"https://www.graphcanon.com/tools/pytorch-serve.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/pytorch-serve","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=pytorch-serve","shared_categories":["inference-serving"]},{"slug":"superlinked-sie","name":"sie","tagline":"Open-source inference server and production cluster for all the models your agent needs.","github_url":"https://github.com/superlinked/sie","owner":"superlinked","repo":"sie","owner_avatar_url":"https://avatars.githubusercontent.com/u/94243920?v=4","primary_language":"Python","stars":3293,"forks":311,"topics":["bge","colbert","data-pipeline","deep-learning","embeddings","inference","inference-server","information-retrieval","llm","ml","mlops","natural-language-processing","nlp","python","reranking","retrieval","retrieval-augmented-generation","semantic-search","splade","vector-search"],"archived":false,"github_pushed_at":"2026-09-19T01:00:02+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/superlinked-sie","markdown_url":"https://www.graphcanon.com/tools/superlinked-sie.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/superlinked-sie","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=superlinked-sie","shared_categories":["inference-serving"]},{"slug":"modelfoxdotdev-modelfox","name":"modelfox","tagline":"ModelFox simplifies machine learning model training and deployment.","github_url":"https://github.com/modelfoxdotdev/modelfox","owner":"modelfoxdotdev","repo":"modelfox","owner_avatar_url":"https://avatars.githubusercontent.com/u/102552266?v=4","primary_language":"Rust","stars":1467,"forks":63,"topics":["automl","developer-tools","elixir","elixir-lang","go","golang","javascript","js","machine-learning","mlops","python","python3","ruby","ruby-on-rails","rust","rust-crate","rust-lang","rust-library","rustlang"],"archived":false,"github_pushed_at":"2024-08-02T17:23:15+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/modelfoxdotdev-modelfox","markdown_url":"https://www.graphcanon.com/tools/modelfoxdotdev-modelfox.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/modelfoxdotdev-modelfox","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=modelfoxdotdev-modelfox","shared_categories":[]},{"slug":"vercel-modelfusion","name":"modelfusion","tagline":"TypeScript library for building AI applications","github_url":"https://github.com/vercel/modelfusion","owner":"vercel","repo":"modelfusion","owner_avatar_url":"https://avatars.githubusercontent.com/u/14985020?v=4","primary_language":"TypeScript","stars":1319,"forks":95,"topics":["ai","artificial-intelligence","chatbot","claude","dall-e","embedding","gpt-3","huggingface","javascript","js","llamacpp","llm","mistral","multi-modal","ollama","openai","stable-diffusion","ts","typescript","whisper"],"archived":true,"github_pushed_at":"2024-07-19T15:17:19+00:00","maintenance_label":"Archived","url":"https://www.graphcanon.com/tools/vercel-modelfusion","markdown_url":"https://www.graphcanon.com/tools/vercel-modelfusion.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/vercel-modelfusion","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=vercel-modelfusion","shared_categories":[]},{"slug":"microsoft-sarathi-serve","name":"sarathi-serve","tagline":"A low-latency and high-throughput serving engine for LLMs","github_url":"https://github.com/microsoft/sarathi-serve","owner":"microsoft","repo":"sarathi-serve","owner_avatar_url":"https://avatars.githubusercontent.com/u/6154722?v=4","primary_language":"Python","stars":527,"forks":67,"topics":["llama","llm-inference","pytorch","transformer"],"archived":false,"github_pushed_at":"2026-01-08T05:10:57+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/microsoft-sarathi-serve","markdown_url":"https://www.graphcanon.com/tools/microsoft-sarathi-serve.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/microsoft-sarathi-serve","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=microsoft-sarathi-serve","shared_categories":["inference-serving"]}]}}