{"data":{"node":{"slug":"tensorflow-serving","name":"serving","tagline":"A flexible, high-performance serving system for machine learning models","github_url":"https://github.com/tensorflow/serving","owner":"tensorflow","repo":"serving","owner_avatar_url":"https://avatars.githubusercontent.com/u/15658638?v=4","primary_language":"C++","stars":6359,"forks":2204,"topics":["cpp","deep-learning","deep-neural-networks","machine-learning","ml","neural-network","python","serving","tensorflow"],"archived":false,"github_pushed_at":"2026-07-30T07:02:43+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/tensorflow-serving","markdown_url":"https://www.graphcanon.com/tools/tensorflow-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/tensorflow-serving","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=tensorflow-serving"},"categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"}],"tags":[{"slug":"cpp","name":"cpp"},{"slug":"deep-learning","name":"deep-learning"},{"slug":"deep-neural-networks","name":"deep-neural-networks"},{"slug":"machine-learning","name":"machine-learning"},{"slug":"ml","name":"ml"},{"slug":"neural-network","name":"neural-network"},{"slug":"python","name":"python"},{"slug":"tensorflow","name":"tensorflow"}],"edges":[{"type":"integrates_with","direction":"in","explanation":"TensorFlow Serving enables high-performance serving of machine learning models with TensorFlow, often used in deployment scenarios that involve models created using TensorFlow as the framework.","successor_context":null,"tool":{"slug":"tensorflow-tensorflow","name":"tensorflow","tagline":"An Open Source Machine Learning Framework for Everyone","github_url":"https://github.com/tensorflow/tensorflow","owner":"tensorflow","repo":"tensorflow","owner_avatar_url":"https://avatars.githubusercontent.com/u/15658638?v=4","primary_language":"C++","stars":196758,"forks":75773,"topics":["deep-learning","deep-neural-networks","distributed","machine-learning","ml","neural-network","python","tensorflow"],"archived":false,"github_pushed_at":"2026-08-03T06:01:06+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/tensorflow-tensorflow","markdown_url":"https://www.graphcanon.com/tools/tensorflow-tensorflow.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/tensorflow-tensorflow","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=tensorflow-tensorflow"}}],"neighbours":[{"slug":"bentoml-bentoml","name":"BentoML","tagline":"The easiest way to serve AI apps and models","github_url":"https://github.com/bentoml/BentoML","owner":"bentoml","repo":"BentoML","owner_avatar_url":"https://avatars.githubusercontent.com/u/49176046?v=4","primary_language":"Python","stars":8793,"forks":1010,"topics":["ai-inference","deep-learning","generative-ai","inference-platform","llm","llm-inference","llm-serving","llmops","machine-learning","ml-engineering","mlops","model-inference-service","model-serving","multimodal","python"],"archived":false,"github_pushed_at":"2026-08-03T17:00:21+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/bentoml-bentoml","markdown_url":"https://www.graphcanon.com/tools/bentoml-bentoml.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/bentoml-bentoml","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=bentoml-bentoml","shared_categories":["inference-serving"]},{"slug":"ai-dynamo-dynamo","name":"dynamo","tagline":"A Datacenter Scale Distributed Inference Serving Framework","github_url":"https://github.com/ai-dynamo/dynamo","owner":"ai-dynamo","repo":"dynamo","owner_avatar_url":"https://avatars.githubusercontent.com/u/201626793?v=4","primary_language":"Rust","stars":7845,"forks":1486,"topics":["diffusion","disaggregated-serving","kubernetes","llm-inference","omni","routing-engine","rust","sglang","tensorrt-llm","vllm"],"archived":false,"github_pushed_at":"2026-08-24T17:59:43+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/ai-dynamo-dynamo","markdown_url":"https://www.graphcanon.com/tools/ai-dynamo-dynamo.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ai-dynamo-dynamo","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ai-dynamo-dynamo","shared_categories":["inference-serving"]},{"slug":"flashinfer-ai-flashinfer","name":"flashinfer","tagline":"FlashInfer is a kernel library for serving large language models","github_url":"https://github.com/flashinfer-ai/flashinfer","owner":"flashinfer-ai","repo":"flashinfer","owner_avatar_url":"https://avatars.githubusercontent.com/u/145061914?v=4","primary_language":"Python","stars":6231,"forks":1327,"topics":["attention","cuda","distributed-inference","gpu","jit","large-large-models","llm-inference","moe","nvidia","pytorch"],"archived":false,"github_pushed_at":"2026-08-24T17:00:11+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/flashinfer-ai-flashinfer","markdown_url":"https://www.graphcanon.com/tools/flashinfer-ai-flashinfer.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/flashinfer-ai-flashinfer","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=flashinfer-ai-flashinfer","shared_categories":["inference-serving"]},{"slug":"kserve-kserve","name":"kserve","tagline":"Standardized Distributed Generative and Predictive AI Inference Platform for Scalable, Multi-Framework Deployment on Kubernetes","github_url":"https://github.com/kserve/kserve","owner":"kserve","repo":"kserve","owner_avatar_url":"https://avatars.githubusercontent.com/u/83512434?v=4","primary_language":"Go","stars":5826,"forks":1632,"topics":["artificial-intelligence","cncf","genai","hacktoberfest","istio","k8s","knative","kserve","kubeflow","kubernetes","llm-inference","machine-learning","mlops","model-interpretability","model-serving","pytorch","service-mesh","tensorflow","vllm","xgboost"],"archived":false,"github_pushed_at":"2026-08-24T02:37:23+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/kserve-kserve","markdown_url":"https://www.graphcanon.com/tools/kserve-kserve.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/kserve-kserve","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=kserve-kserve","shared_categories":["inference-serving"]},{"slug":"seldonio-seldon-core","name":"seldon-core","tagline":"An MLOps framework to package, deploy, monitor and manage thousands of production machine learning models","github_url":"https://github.com/SeldonIO/seldon-core","owner":"SeldonIO","repo":"seldon-core","owner_avatar_url":"https://avatars.githubusercontent.com/u/10297834?v=4","primary_language":"Go","stars":4765,"forks":867,"topics":["aiops","deployment","kubernetes","machine-learning","machine-learning-operations","mlops","production-machine-learning","serving"],"archived":false,"github_pushed_at":"2026-03-23T11:39:54+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/seldonio-seldon-core","markdown_url":"https://www.graphcanon.com/tools/seldonio-seldon-core.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/seldonio-seldon-core","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=seldonio-seldon-core","shared_categories":["inference-serving"]},{"slug":"pytorch-serve","name":"serve","tagline":"Serve, optimize and scale PyTorch models in production","github_url":"https://github.com/pytorch/serve","owner":"pytorch","repo":"serve","owner_avatar_url":"https://avatars.githubusercontent.com/u/21003710?v=4","primary_language":"Java","stars":4350,"forks":882,"topics":["cpu","deep-learning","docker","gpu","kubernetes","machine-learning","metrics","mlops","optimization","pytorch","serving"],"archived":true,"github_pushed_at":"2025-08-06T19:17:08+00:00","maintenance_label":"Archived","url":"https://www.graphcanon.com/tools/pytorch-serve","markdown_url":"https://www.graphcanon.com/tools/pytorch-serve.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/pytorch-serve","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=pytorch-serve","shared_categories":["inference-serving"]},{"slug":"michaelfeil-infinity","name":"infinity","tagline":"High-throughput, low-latency serving engine for text-embeddings and various models","github_url":"https://github.com/michaelfeil/infinity","owner":"michaelfeil","repo":"infinity","owner_avatar_url":"https://avatars.githubusercontent.com/u/63565275?v=4","primary_language":"Python","stars":2907,"forks":196,"topics":["bert-embeddings","llm","text-embeddings"],"archived":false,"github_pushed_at":"2026-03-24T03:59:47+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/michaelfeil-infinity","markdown_url":"https://www.graphcanon.com/tools/michaelfeil-infinity.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/michaelfeil-infinity","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=michaelfeil-infinity","shared_categories":["inference-serving"]},{"slug":"superlinked-sie","name":"sie","tagline":"Open-source inference server and production cluster for all the models your agent needs.","github_url":"https://github.com/superlinked/sie","owner":"superlinked","repo":"sie","owner_avatar_url":"https://avatars.githubusercontent.com/u/94243920?v=4","primary_language":"Python","stars":2804,"forks":272,"topics":["bge","colbert","data-pipeline","deep-learning","embeddings","inference","inference-server","information-retrieval","llm","ml","mlops","natural-language-processing","nlp","python","reranking","retrieval","retrieval-augmented-generation","semantic-search","splade","vector-search"],"archived":false,"github_pushed_at":"2026-08-21T20:28:04+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/superlinked-sie","markdown_url":"https://www.graphcanon.com/tools/superlinked-sie.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/superlinked-sie","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=superlinked-sie","shared_categories":["inference-serving"]},{"slug":"microsoft-pai","name":"pai","tagline":"Resource scheduling and cluster management for AI","github_url":"https://github.com/microsoft/pai","owner":"microsoft","repo":"pai","owner_avatar_url":"https://avatars.githubusercontent.com/u/6154722?v=4","primary_language":"JavaScript","stars":2686,"forks":549,"topics":["ai","artificial-intelligence","chainer","cloud","cluster-management","cluster-manager","gpu","gpu-cluster","gpu-computing","gpu-scheduler","jupyter","kubernetes","machine-learning","model-training","on-premise","pytorch","resource-management","scheduling","tensorflow"],"archived":true,"github_pushed_at":"2024-06-06T07:56:07+00:00","maintenance_label":"Archived","url":"https://www.graphcanon.com/tools/microsoft-pai","markdown_url":"https://www.graphcanon.com/tools/microsoft-pai.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/microsoft-pai","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=microsoft-pai","shared_categories":["inference-serving"]},{"slug":"vertaai-modeldb","name":"modeldb","tagline":"Open Source ML Model Versioning Metadata and Experiment Management","github_url":"https://github.com/VertaAI/modeldb","owner":"VertaAI","repo":"modeldb","owner_avatar_url":"https://avatars.githubusercontent.com/u/45020510?v=4","primary_language":"Java","stars":1749,"forks":289,"topics":["machine-learning","mit","model-management","model-versioning","modeldb","verta"],"archived":false,"github_pushed_at":"2024-07-23T17:06:34+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/vertaai-modeldb","markdown_url":"https://www.graphcanon.com/tools/vertaai-modeldb.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/vertaai-modeldb","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=vertaai-modeldb","shared_categories":[]},{"slug":"ebhy-budgetml","name":"budgetml","tagline":"Deploys ML inference service economically","github_url":"https://github.com/ebhy/budgetml","owner":"ebhy","repo":"budgetml","owner_avatar_url":"https://avatars.githubusercontent.com/u/76654256?v=4","primary_language":"Python","stars":1343,"forks":65,"topics":["api","data-science","deployment","fastapi","inference","machine-learning","mlops"],"archived":false,"github_pushed_at":"2024-02-12T17:29:24+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/ebhy-budgetml","markdown_url":"https://www.graphcanon.com/tools/ebhy-budgetml.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ebhy-budgetml","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ebhy-budgetml","shared_categories":["inference-serving"]},{"slug":"kubeai-project-kubeai","name":"kubeai","tagline":"AI Inference Operator for Kubernetes","github_url":"https://github.com/kubeai-project/kubeai","owner":"kubeai-project","repo":"kubeai","owner_avatar_url":"https://avatars.githubusercontent.com/u/232319222?v=4","primary_language":"Go","stars":1237,"forks":131,"topics":["ai","autoscaler","faster-whisper","inference-operator","k8s","kubernetes","llm","ollama","ollama-operator","openai-api","vllm","vllm-operator","whisper"],"archived":false,"github_pushed_at":"2026-07-31T01:04:47+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/kubeai-project-kubeai","markdown_url":"https://www.graphcanon.com/tools/kubeai-project-kubeai.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/kubeai-project-kubeai","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=kubeai-project-kubeai","shared_categories":["inference-serving"]}]}}