{"data":{"slug":"tensorflow-serving","name":"serving","tagline":"A flexible, high-performance serving system for machine learning models","github_url":"https://github.com/tensorflow/serving","owner":"tensorflow","repo":"serving","owner_avatar_url":"https://avatars.githubusercontent.com/u/15658638?v=4","primary_language":"C++","stars":6359,"forks":2204,"topics":["cpp","deep-learning","deep-neural-networks","machine-learning","ml","neural-network","python","serving","tensorflow"],"archived":false,"github_pushed_at":"2026-07-30T07:02:43+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/tensorflow-serving","markdown_url":"https://www.graphcanon.com/tools/tensorflow-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/tensorflow-serving","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=tensorflow-serving","description":"A flexible, high-performance serving system for machine learning models","homepage_url":"https://www.tensorflow.org/serving","license":"Apache-2.0","open_issues":95,"watchers":217,"ai_summary":"TensorFlow Serving is designed to serve machine learning models with low latency and high throughput.","readme_excerpt":"# Download the TensorFlow Serving Docker image and repo\ndocker pull tensorflow/serving\n\ngit clone https://github.com/tensorflow/serving","github_created_at":"2016-01-26T21:48:20+00:00","created_at":"2026-07-11T23:12:33.413454+00:00","updated_at":"2026-08-02T06:00:07.91277+00:00","categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"}],"tags":[{"slug":"cpp","name":"cpp"},{"slug":"deep-learning","name":"deep-learning"},{"slug":"deep-neural-networks","name":"deep-neural-networks"},{"slug":"machine-learning","name":"machine-learning"},{"slug":"ml","name":"ml"},{"slug":"neural-network","name":"neural-network"},{"slug":"python","name":"python"},{"slug":"tensorflow","name":"tensorflow"}],"trust":{"provenance":{"is_fork":false,"github_id":50461701,"owner_type":"Organization","methodology":"github_public_v1","parent_repo":null,"near_duplicate_slugs":[]},"computed_at":"2026-08-02T06:00:07.203Z","maintenance":{"label":"Very active","score":96,"methodology":"github_public_v1","releases_90d":1,"days_since_push":2,"last_release_at":"2026-06-02T18:48:59Z"},"security_summary":{"status":"no_lockfile","scanner":null,"low_count":0,"high_count":0,"last_scan_at":"2026-07-11T23:12:42.769Z","medium_count":0,"scan_profile":"none","critical_count":0}},"capability_facts":{"scan":{"source":"repo_scan","observed_at":"2026-08-02T06:00:07.638Z"},"languages":{"value":["c++"],"source":"github.language","observed_at":"2026-08-02T06:00:07.638Z"},"license_spdx":{"value":"Apache-2.0","source":"github.license","observed_at":"2026-08-02T06:00:07.638Z"}},"decision_facts":{"hosting":null,"pricing":null,"requirements":null,"constraints":null,"when_to_use":["When you have an existing TensorFlow model that benefits from ultra-low latency and high throughput, especially in production environments where performance is critical.","If your application is written in C++ or Python and requires tightly integrated, fast inference capabilities.","For large-scale deployments where maintaining low service latency under varying load conditions is essential."],"when_not_to_use":["When working with smaller models that don't require the scalability features of TensorFlow Serving; simpler serving solutions like Flask servers might suffice.","If your model development and deployment stack does not include TensorFlow, finding it easier to stick with libraries specific to your existing framework (like PyTorch's TorchServe).","In cases where flexibility in customizing the serving environment is more critical than out-of-the-box performance benefits."],"source":"enrich:decision_facts","observed_at":"2026-07-12T15:47:27.181Z"},"constraint_facets":null,"decision_summary":[{"label":"Adopt for","value":"TensorFlow Serving is a high-performance machine learning serving system built for low latency and high throughput scenarios."}]}}