{"data":{"node":{"slug":"autodeployai-ai-serving","name":"ai-serving","tagline":"Serving AI/ML models in open standard formats PMML and ONNX with HTTP and gRPC endpoints","github_url":"https://github.com/autodeployai/ai-serving","owner":"autodeployai","repo":"ai-serving","owner_avatar_url":"https://avatars.githubusercontent.com/u/50666721?v=4","primary_language":"Scala","stars":166,"forks":31,"topics":["ai-serving","inference","inference-server","onnx","onnx-grpc","onnx-inference","onnx-models","onnx-realtime","onnx-rest","pmml","pmml-deployment","pmml-grpc","pmml-inference","pmml-model","pmml-realtime","pmml-rest"],"archived":false,"github_pushed_at":"2026-02-24T02:56:29+00:00","maintenance_label":"Slowing","stars_delta_30d":0,"url":"https://www.graphcanon.com/tools/autodeployai-ai-serving","markdown_url":"https://www.graphcanon.com/tools/autodeployai-ai-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/autodeployai-ai-serving","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=autodeployai-ai-serving"},"categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"}],"tags":[{"slug":"ai-serving","name":"ai-serving"},{"slug":"grpc","name":"grpc"},{"slug":"inference-server","name":"inference-server"},{"slug":"onnx","name":"onnx"},{"slug":"pmml-deployment","name":"pmml-deployment"},{"slug":"rest-api","name":"rest-api"}],"edges":[],"neighbours":[{"slug":"jundot-omlx","name":"omlx","tagline":"LLM inference server with continuous batching and SSD caching for Apple Silicon","github_url":"https://github.com/jundot/omlx","owner":"jundot","repo":"omlx","owner_avatar_url":"https://avatars.githubusercontent.com/u/64250138?v=4","primary_language":"Python","stars":21934,"forks":1899,"topics":["apple-silicon","inference-server","llm","macos","mlx","openai-api"],"archived":false,"github_pushed_at":"2026-09-20T02:23:26+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/jundot-omlx","markdown_url":"https://www.graphcanon.com/tools/jundot-omlx.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/jundot-omlx","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=jundot-omlx","shared_categories":["inference-serving"]},{"slug":"xorbitsai-inference","name":"inference","tagline":"Unified production-ready inference API for various LLMs and models","github_url":"https://github.com/xorbitsai/inference","owner":"xorbitsai","repo":"inference","owner_avatar_url":"https://avatars.githubusercontent.com/u/109655068?v=4","primary_language":"Python","stars":9575,"forks":866,"topics":["artificial-intelligence","deployment","diffusers","gemma","glm","glm-5-3","inference","kimi","kimi-k3","llama-cpp","llamacpp","llm","machine-learning","openai-api","pytorch","qwen","sglang","transformers","vllm","whisper"],"archived":false,"github_pushed_at":"2026-09-18T06:04:24+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/xorbitsai-inference","markdown_url":"https://www.graphcanon.com/tools/xorbitsai-inference.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/xorbitsai-inference","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=xorbitsai-inference","shared_categories":["inference-serving"]},{"slug":"bentoml-bentoml","name":"BentoML","tagline":"The easiest way to serve AI apps and models","github_url":"https://github.com/bentoml/BentoML","owner":"bentoml","repo":"BentoML","owner_avatar_url":"https://avatars.githubusercontent.com/u/49176046?v=4","primary_language":"Python","stars":8847,"forks":1032,"topics":["ai-inference","deep-learning","generative-ai","inference-platform","llm","llm-inference","llm-serving","llmops","machine-learning","ml-engineering","mlops","model-inference-service","model-serving","multimodal","python"],"archived":false,"github_pushed_at":"2026-09-07T17:43:01+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/bentoml-bentoml","markdown_url":"https://www.graphcanon.com/tools/bentoml-bentoml.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/bentoml-bentoml","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=bentoml-bentoml","shared_categories":["inference-serving"]},{"slug":"kserve-kserve","name":"kserve","tagline":"Standardized Distributed Generative and Predictive AI Inference Platform for Scalable, Multi-Framework Deployment on Kubernetes","github_url":"https://github.com/kserve/kserve","owner":"kserve","repo":"kserve","owner_avatar_url":"https://avatars.githubusercontent.com/u/83512434?v=4","primary_language":"Go","stars":5970,"forks":1689,"topics":["artificial-intelligence","cncf","genai","hacktoberfest","istio","k8s","knative","kserve","kubeflow","kubernetes","llm-inference","machine-learning","mlops","model-interpretability","model-serving","pytorch","service-mesh","tensorflow","vllm","xgboost"],"archived":false,"github_pushed_at":"2026-09-19T13:08:33+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/kserve-kserve","markdown_url":"https://www.graphcanon.com/tools/kserve-kserve.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/kserve-kserve","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=kserve-kserve","shared_categories":["inference-serving"]},{"slug":"seldonio-seldon-core","name":"seldon-core","tagline":"An MLOps framework to package, deploy, monitor and manage thousands of production machine learning models","github_url":"https://github.com/SeldonIO/seldon-core","owner":"SeldonIO","repo":"seldon-core","owner_avatar_url":"https://avatars.githubusercontent.com/u/10297834?v=4","primary_language":"Go","stars":4780,"forks":868,"topics":["aiops","deployment","kubernetes","machine-learning","machine-learning-operations","mlops","production-machine-learning","serving"],"archived":false,"github_pushed_at":"2026-03-23T11:39:54+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/seldonio-seldon-core","markdown_url":"https://www.graphcanon.com/tools/seldonio-seldon-core.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/seldonio-seldon-core","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=seldonio-seldon-core","shared_categories":["inference-serving"]},{"slug":"pytorch-serve","name":"serve","tagline":"Serve, optimize and scale PyTorch models in production","github_url":"https://github.com/pytorch/serve","owner":"pytorch","repo":"serve","owner_avatar_url":"https://avatars.githubusercontent.com/u/21003710?v=4","primary_language":"Java","stars":4344,"forks":880,"topics":["cpu","deep-learning","docker","gpu","kubernetes","machine-learning","metrics","mlops","optimization","pytorch","serving"],"archived":true,"github_pushed_at":"2025-08-06T19:17:08+00:00","maintenance_label":"Archived","url":"https://www.graphcanon.com/tools/pytorch-serve","markdown_url":"https://www.graphcanon.com/tools/pytorch-serve.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/pytorch-serve","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=pytorch-serve","shared_categories":["inference-serving"]},{"slug":"superlinked-sie","name":"sie","tagline":"Open-source inference server and production cluster for all the models your agent needs.","github_url":"https://github.com/superlinked/sie","owner":"superlinked","repo":"sie","owner_avatar_url":"https://avatars.githubusercontent.com/u/94243920?v=4","primary_language":"Python","stars":3293,"forks":311,"topics":["bge","colbert","data-pipeline","deep-learning","embeddings","inference","inference-server","information-retrieval","llm","ml","mlops","natural-language-processing","nlp","python","reranking","retrieval","retrieval-augmented-generation","semantic-search","splade","vector-search"],"archived":false,"github_pushed_at":"2026-09-19T01:00:02+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/superlinked-sie","markdown_url":"https://www.graphcanon.com/tools/superlinked-sie.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/superlinked-sie","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=superlinked-sie","shared_categories":["inference-serving"]},{"slug":"b4rtaz-distributed-llama","name":"distributed-llama","tagline":"Distributed LLM inference using home devices cluster","github_url":"https://github.com/b4rtaz/distributed-llama","owner":"b4rtaz","repo":"distributed-llama","owner_avatar_url":"https://avatars.githubusercontent.com/u/12797776?v=4","primary_language":"C++","stars":3060,"forks":250,"topics":["distributed-computing","distributed-llm","llama2","llama3","llm","llm-inference","llms","neural-network","open-llm"],"archived":false,"github_pushed_at":"2026-07-05T16:47:20+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/b4rtaz-distributed-llama","markdown_url":"https://www.graphcanon.com/tools/b4rtaz-distributed-llama.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/b4rtaz-distributed-llama","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=b4rtaz-distributed-llama","shared_categories":["inference-serving"]},{"slug":"ailia-ai-ailia-models","name":"ailia-models","tagline":"Repository of pre-trained AI models for ailia SDK","github_url":"https://github.com/ailia-ai/ailia-models","owner":"ailia-ai","repo":"ailia-models","owner_avatar_url":"https://avatars.githubusercontent.com/u/55011924?v=4","primary_language":"Python","stars":2392,"forks":365,"topics":["action-recognition","anomaly-detection","audio-processing","background-removal","crowd-counting","deep-learning","embeddings","face-detection","face-recognition","fashion-ai","gan","hand-detection","image-classification","image-segmentation","llm","neural-network","object-detection","object-recognition","object-tracking","pose-estimation"],"archived":false,"github_pushed_at":"2026-09-16T08:01:06+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/ailia-ai-ailia-models","markdown_url":"https://www.graphcanon.com/tools/ailia-ai-ailia-models.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ailia-ai-ailia-models","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ailia-ai-ailia-models","shared_categories":[]},{"slug":"minimaxir-automl-gs","name":"automl-gs","tagline":"Automatically generate machine-learning models and code with input CSV and target field","github_url":"https://github.com/minimaxir/automl-gs","owner":"minimaxir","repo":"automl-gs","owner_avatar_url":"https://avatars.githubusercontent.com/u/2179708?v=4","primary_language":"Python","stars":1866,"forks":180,"topics":["automl","keras","machine-learning","python","tensorflow","xgboost"],"archived":false,"github_pushed_at":"2019-10-22T11:20:40+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/minimaxir-automl-gs","markdown_url":"https://www.graphcanon.com/tools/minimaxir-automl-gs.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/minimaxir-automl-gs","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=minimaxir-automl-gs","shared_categories":[]},{"slug":"waybarrios-vllm-mlx","name":"vllm-mlx","tagline":"Server for LLMs and vision-language models compatible with Apple Silicon","github_url":"https://github.com/waybarrios/vllm-mlx","owner":"waybarrios","repo":"vllm-mlx","owner_avatar_url":"https://avatars.githubusercontent.com/u/6794828?v=4","primary_language":"Python","stars":1588,"forks":222,"topics":["anthropic","anthropic-api","apple-silicon","claude-code","continuous-batching","inference-server","llm","local-llm","macos","mcp","mlx","multimodal-ai","openai","openai-api","openai-compatible","speech-to-text","text-to-speech","tool-calling","vision-language-model","vllm"],"archived":false,"github_pushed_at":"2026-09-19T18:11:07+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/waybarrios-vllm-mlx","markdown_url":"https://www.graphcanon.com/tools/waybarrios-vllm-mlx.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/waybarrios-vllm-mlx","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=waybarrios-vllm-mlx","shared_categories":["inference-serving"]},{"slug":"modelfoxdotdev-modelfox","name":"modelfox","tagline":"ModelFox simplifies machine learning model training and deployment.","github_url":"https://github.com/modelfoxdotdev/modelfox","owner":"modelfoxdotdev","repo":"modelfox","owner_avatar_url":"https://avatars.githubusercontent.com/u/102552266?v=4","primary_language":"Rust","stars":1467,"forks":63,"topics":["automl","developer-tools","elixir","elixir-lang","go","golang","javascript","js","machine-learning","mlops","python","python3","ruby","ruby-on-rails","rust","rust-crate","rust-lang","rust-library","rustlang"],"archived":false,"github_pushed_at":"2024-08-02T17:23:15+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/modelfoxdotdev-modelfox","markdown_url":"https://www.graphcanon.com/tools/modelfoxdotdev-modelfox.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/modelfoxdotdev-modelfox","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=modelfoxdotdev-modelfox","shared_categories":[]}]}}