{"data":{"node":{"slug":"kserve-kserve","name":"kserve","tagline":"Standardized Distributed Generative and Predictive AI Inference Platform for Scalable, Multi-Framework Deployment on Kubernetes","github_url":"https://github.com/kserve/kserve","owner":"kserve","repo":"kserve","owner_avatar_url":"https://avatars.githubusercontent.com/u/83512434?v=4","primary_language":"Go","stars":5826,"forks":1632,"topics":["artificial-intelligence","cncf","genai","hacktoberfest","istio","k8s","knative","kserve","kubeflow","kubernetes","llm-inference","machine-learning","mlops","model-interpretability","model-serving","pytorch","service-mesh","tensorflow","vllm","xgboost"],"archived":false,"github_pushed_at":"2026-08-24T02:37:23+00:00","maintenance_label":"Very active","stars_delta_30d":95,"url":"https://www.graphcanon.com/tools/kserve-kserve","markdown_url":"https://www.graphcanon.com/tools/kserve-kserve.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/kserve-kserve","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=kserve-kserve"},"categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"}],"tags":[{"slug":"artificial-intelligence","name":"artificial-intelligence"},{"slug":"cncf","name":"cncf"},{"slug":"genai","name":"genai"},{"slug":"hacktoberfest","name":"hacktoberfest"},{"slug":"istio","name":"istio"},{"slug":"k8s","name":"k8s"},{"slug":"knative","name":"knative"},{"slug":"kserve","name":"kserve"}],"edges":[],"neighbours":[{"slug":"jina-ai-serve","name":"serve","tagline":"Build multimodal AI applications with cloud-native stack","github_url":"https://github.com/jina-ai/serve","owner":"jina-ai","repo":"serve","owner_avatar_url":"https://avatars.githubusercontent.com/u/60539444?v=4","primary_language":"Python","stars":21863,"forks":2243,"topics":["cloud-native","cncf","deep-learning","docker","fastapi","framework","generative-ai","grpc","jaeger","kubernetes","llmops","machine-learning","microservice","mlops","multimodal","neural-search","opentelemetry","orchestration","pipeline","prometheus"],"archived":false,"github_pushed_at":"2025-03-24T13:59:54+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/jina-ai-serve","markdown_url":"https://www.graphcanon.com/tools/jina-ai-serve.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/jina-ai-serve","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=jina-ai-serve","shared_categories":["inference-serving"]},{"slug":"skypilot-org-skypilot","name":"skypilot","tagline":"Run, manage, and scale AI workloads on any AI infrastructure.","github_url":"https://github.com/skypilot-org/skypilot","owner":"skypilot-org","repo":"skypilot","owner_avatar_url":"https://avatars.githubusercontent.com/u/109387420?v=4","primary_language":"Python","stars":10456,"forks":1175,"topics":["cloud-computing","cloud-management","cost-optimization","deep-learning","distributed-training","gpu","hyperparameter-tuning","job-queue","job-scheduler","llm-serving","llm-training","machine-learning","ml-infrastructure","ml-platform","mlops","multicloud","slurm","spot-instances","tpu"],"archived":false,"github_pushed_at":"2026-08-07T10:25:54+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/skypilot-org-skypilot","markdown_url":"https://www.graphcanon.com/tools/skypilot-org-skypilot.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/skypilot-org-skypilot","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=skypilot-org-skypilot","shared_categories":["inference-serving"]},{"slug":"bentoml-bentoml","name":"BentoML","tagline":"The easiest way to serve AI apps and models","github_url":"https://github.com/bentoml/BentoML","owner":"bentoml","repo":"BentoML","owner_avatar_url":"https://avatars.githubusercontent.com/u/49176046?v=4","primary_language":"Python","stars":8793,"forks":1010,"topics":["ai-inference","deep-learning","generative-ai","inference-platform","llm","llm-inference","llm-serving","llmops","machine-learning","ml-engineering","mlops","model-inference-service","model-serving","multimodal","python"],"archived":false,"github_pushed_at":"2026-08-03T17:00:21+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/bentoml-bentoml","markdown_url":"https://www.graphcanon.com/tools/bentoml-bentoml.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/bentoml-bentoml","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=bentoml-bentoml","shared_categories":["inference-serving"]},{"slug":"ai-dynamo-dynamo","name":"dynamo","tagline":"A Datacenter Scale Distributed Inference Serving Framework","github_url":"https://github.com/ai-dynamo/dynamo","owner":"ai-dynamo","repo":"dynamo","owner_avatar_url":"https://avatars.githubusercontent.com/u/201626793?v=4","primary_language":"Rust","stars":7845,"forks":1486,"topics":["diffusion","disaggregated-serving","kubernetes","llm-inference","omni","routing-engine","rust","sglang","tensorrt-llm","vllm"],"archived":false,"github_pushed_at":"2026-08-24T17:59:43+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/ai-dynamo-dynamo","markdown_url":"https://www.graphcanon.com/tools/ai-dynamo-dynamo.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ai-dynamo-dynamo","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ai-dynamo-dynamo","shared_categories":["inference-serving"]},{"slug":"tensorflow-serving","name":"serving","tagline":"A flexible, high-performance serving system for machine learning models","github_url":"https://github.com/tensorflow/serving","owner":"tensorflow","repo":"serving","owner_avatar_url":"https://avatars.githubusercontent.com/u/15658638?v=4","primary_language":"C++","stars":6359,"forks":2204,"topics":["cpp","deep-learning","deep-neural-networks","machine-learning","ml","neural-network","python","serving","tensorflow"],"archived":false,"github_pushed_at":"2026-07-30T07:02:43+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/tensorflow-serving","markdown_url":"https://www.graphcanon.com/tools/tensorflow-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/tensorflow-serving","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=tensorflow-serving","shared_categories":["inference-serving"]},{"slug":"flashinfer-ai-flashinfer","name":"flashinfer","tagline":"FlashInfer is a kernel library for serving large language models","github_url":"https://github.com/flashinfer-ai/flashinfer","owner":"flashinfer-ai","repo":"flashinfer","owner_avatar_url":"https://avatars.githubusercontent.com/u/145061914?v=4","primary_language":"Python","stars":6231,"forks":1327,"topics":["attention","cuda","distributed-inference","gpu","jit","large-large-models","llm-inference","moe","nvidia","pytorch"],"archived":false,"github_pushed_at":"2026-08-24T17:00:11+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/flashinfer-ai-flashinfer","markdown_url":"https://www.graphcanon.com/tools/flashinfer-ai-flashinfer.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/flashinfer-ai-flashinfer","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=flashinfer-ai-flashinfer","shared_categories":["inference-serving"]},{"slug":"seldonio-seldon-core","name":"seldon-core","tagline":"An MLOps framework to package, deploy, monitor and manage thousands of production machine learning models","github_url":"https://github.com/SeldonIO/seldon-core","owner":"SeldonIO","repo":"seldon-core","owner_avatar_url":"https://avatars.githubusercontent.com/u/10297834?v=4","primary_language":"Go","stars":4765,"forks":867,"topics":["aiops","deployment","kubernetes","machine-learning","machine-learning-operations","mlops","production-machine-learning","serving"],"archived":false,"github_pushed_at":"2026-03-23T11:39:54+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/seldonio-seldon-core","markdown_url":"https://www.graphcanon.com/tools/seldonio-seldon-core.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/seldonio-seldon-core","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=seldonio-seldon-core","shared_categories":["inference-serving"]},{"slug":"kubeflow-pipelines","name":"pipelines","tagline":"Machine Learning Pipelines for Kubeflow","github_url":"https://github.com/kubeflow/pipelines","owner":"kubeflow","repo":"pipelines","owner_avatar_url":"https://avatars.githubusercontent.com/u/33164907?v=4","primary_language":"Python","stars":4173,"forks":2075,"topics":["data-science","kubeflow","kubeflow-pipelines","kubernetes","machine-learning","mlops","pipeline"],"archived":false,"github_pushed_at":"2026-08-03T15:35:16+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/kubeflow-pipelines","markdown_url":"https://www.graphcanon.com/tools/kubeflow-pipelines.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/kubeflow-pipelines","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=kubeflow-pipelines","shared_categories":["inference-serving"]},{"slug":"superlinked-sie","name":"sie","tagline":"Open-source inference server and production cluster for all the models your agent needs.","github_url":"https://github.com/superlinked/sie","owner":"superlinked","repo":"sie","owner_avatar_url":"https://avatars.githubusercontent.com/u/94243920?v=4","primary_language":"Python","stars":2804,"forks":272,"topics":["bge","colbert","data-pipeline","deep-learning","embeddings","inference","inference-server","information-retrieval","llm","ml","mlops","natural-language-processing","nlp","python","reranking","retrieval","retrieval-augmented-generation","semantic-search","splade","vector-search"],"archived":false,"github_pushed_at":"2026-08-21T20:28:04+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/superlinked-sie","markdown_url":"https://www.graphcanon.com/tools/superlinked-sie.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/superlinked-sie","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=superlinked-sie","shared_categories":["inference-serving"]},{"slug":"microsoft-pai","name":"pai","tagline":"Resource scheduling and cluster management for AI","github_url":"https://github.com/microsoft/pai","owner":"microsoft","repo":"pai","owner_avatar_url":"https://avatars.githubusercontent.com/u/6154722?v=4","primary_language":"JavaScript","stars":2686,"forks":549,"topics":["ai","artificial-intelligence","chainer","cloud","cluster-management","cluster-manager","gpu","gpu-cluster","gpu-computing","gpu-scheduler","jupyter","kubernetes","machine-learning","model-training","on-premise","pytorch","resource-management","scheduling","tensorflow"],"archived":true,"github_pushed_at":"2024-06-06T07:56:07+00:00","maintenance_label":"Archived","url":"https://www.graphcanon.com/tools/microsoft-pai","markdown_url":"https://www.graphcanon.com/tools/microsoft-pai.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/microsoft-pai","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=microsoft-pai","shared_categories":["inference-serving"]},{"slug":"dstackai-dstack","name":"dstack","tagline":"Vendor-agnostic orchestration for AI workloads","github_url":"https://github.com/dstackai/dstack","owner":"dstackai","repo":"dstack","owner_avatar_url":"https://avatars.githubusercontent.com/u/54146142?v=4","primary_language":"Python","stars":2219,"forks":250,"topics":["agent-skills","agentic-orchestration","amd","cloud","containers","docker","fine-tuning","gpu","inference","k8s","kubernetes","llms","machine-learning","nvidia","orchestration","python","slurm","training"],"archived":false,"github_pushed_at":"2026-08-23T13:41:13+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/dstackai-dstack","markdown_url":"https://www.graphcanon.com/tools/dstackai-dstack.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/dstackai-dstack","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=dstackai-dstack","shared_categories":["inference-serving"]},{"slug":"kubeflow-trainer","name":"trainer","tagline":"Distributed AI Model Training and LLM Fine-Tuning on Kubernetes","github_url":"https://github.com/kubeflow/trainer","owner":"kubeflow","repo":"trainer","owner_avatar_url":"https://avatars.githubusercontent.com/u/33164907?v=4","primary_language":"Go","stars":2196,"forks":1030,"topics":["ai","distributed","fine-tuning","gpu","huggingface","jax","kubeflow","kubernetes","llm","machine-learning","mlops","python","pytorch","tensorflow","xgboost"],"archived":false,"github_pushed_at":"2026-08-22T02:27:28+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/kubeflow-trainer","markdown_url":"https://www.graphcanon.com/tools/kubeflow-trainer.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/kubeflow-trainer","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=kubeflow-trainer","shared_categories":[]}]}}