{"data":{"node":{"slug":"onnx-onnx-mlir","name":"onnx-mlir","tagline":"ONNX model compiler technology lowering ONNX graphs to MLIR and LLVM bytecodes","github_url":"https://github.com/onnx/onnx-mlir","owner":"onnx","repo":"onnx-mlir","owner_avatar_url":"https://avatars.githubusercontent.com/u/31675368?v=4","primary_language":"C++","stars":1039,"forks":447,"topics":[],"archived":false,"github_pushed_at":"2026-07-31T03:15:39+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/onnx-onnx-mlir","markdown_url":"https://www.graphcanon.com/tools/onnx-onnx-mlir.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/onnx-onnx-mlir","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=onnx-onnx-mlir"},"categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"},{"slug":"model-training","name":"Model Training","url":"https://www.graphcanon.com/categories/model-training","markdown_url":"https://www.graphcanon.com/categories/model-training.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/model-training"}],"tags":[{"slug":"compiler","name":"compiler"},{"slug":"llvm","name":"llvm"},{"slug":"mlir","name":"mlir"},{"slug":"onnx","name":"onnx"},{"slug":"runtime-environments","name":"runtime environments"}],"edges":[],"neighbours":[{"slug":"tensorflow-tensorflow","name":"tensorflow","tagline":"An Open Source Machine Learning Framework for Everyone","github_url":"https://github.com/tensorflow/tensorflow","owner":"tensorflow","repo":"tensorflow","owner_avatar_url":"https://avatars.githubusercontent.com/u/15658638?v=4","primary_language":"C++","stars":196758,"forks":75773,"topics":["deep-learning","deep-neural-networks","distributed","machine-learning","ml","neural-network","python","tensorflow"],"archived":false,"github_pushed_at":"2026-08-03T06:01:06+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/tensorflow-tensorflow","markdown_url":"https://www.graphcanon.com/tools/tensorflow-tensorflow.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/tensorflow-tensorflow","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=tensorflow-tensorflow","shared_categories":["model-training"]},{"slug":"k2-fsa-sherpa-onnx","name":"sherpa-onnx","tagline":"Speech-to-text and related audio processing tools using ONNX with cross-platform support","github_url":"https://github.com/k2-fsa/sherpa-onnx","owner":"k2-fsa","repo":"sherpa-onnx","owner_avatar_url":"https://avatars.githubusercontent.com/u/71431748?v=4","primary_language":"C++","stars":13842,"forks":1596,"topics":["aarch64","android","arm32","asr","cpp","csharp","dotnet","ios","lazarus","linux","macos","mfc","object-pascal","onnx","raspberry-pi","risc-v","speech-to-text","text-to-speech","vits","windows"],"archived":false,"github_pushed_at":"2026-07-29T03:35:02+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/k2-fsa-sherpa-onnx","markdown_url":"https://www.graphcanon.com/tools/k2-fsa-sherpa-onnx.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/k2-fsa-sherpa-onnx","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=k2-fsa-sherpa-onnx","shared_categories":[]},{"slug":"bitsandbytes-foundation-bitsandbytes","name":"bitsandbytes","tagline":"Large language model quantization toolkit for PyTorch.","github_url":"https://github.com/bitsandbytes-foundation/bitsandbytes","owner":"bitsandbytes-foundation","repo":"bitsandbytes","owner_avatar_url":"https://avatars.githubusercontent.com/u/175231607?v=4","primary_language":"Python","stars":8385,"forks":900,"topics":["llm","machine-learning","pytorch","qlora","quantization"],"archived":false,"github_pushed_at":"2026-07-29T18:27:51+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/bitsandbytes-foundation-bitsandbytes","markdown_url":"https://www.graphcanon.com/tools/bitsandbytes-foundation-bitsandbytes.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/bitsandbytes-foundation-bitsandbytes","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=bitsandbytes-foundation-bitsandbytes","shared_categories":["inference-serving"]},{"slug":"linkedin-liger-kernel","name":"Liger-Kernel","tagline":"Efficient Triton Kernels for LLM Training","github_url":"https://github.com/linkedin/Liger-Kernel","owner":"linkedin","repo":"Liger-Kernel","owner_avatar_url":"https://avatars.githubusercontent.com/u/357098?v=4","primary_language":"Python","stars":6555,"forks":573,"topics":["finetuning","gemma2","hacktoberfest","llama","llama3","llm-training","llms","mistral","phi3","triton","triton-kernels"],"archived":false,"github_pushed_at":"2026-08-07T08:48:09+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/linkedin-liger-kernel","markdown_url":"https://www.graphcanon.com/tools/linkedin-liger-kernel.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/linkedin-liger-kernel","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=linkedin-liger-kernel","shared_categories":["model-training"]},{"slug":"tensorflow-serving","name":"serving","tagline":"A flexible, high-performance serving system for machine learning models","github_url":"https://github.com/tensorflow/serving","owner":"tensorflow","repo":"serving","owner_avatar_url":"https://avatars.githubusercontent.com/u/15658638?v=4","primary_language":"C++","stars":6359,"forks":2204,"topics":["cpp","deep-learning","deep-neural-networks","machine-learning","ml","neural-network","python","serving","tensorflow"],"archived":false,"github_pushed_at":"2026-07-30T07:02:43+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/tensorflow-serving","markdown_url":"https://www.graphcanon.com/tools/tensorflow-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/tensorflow-serving","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=tensorflow-serving","shared_categories":["inference-serving"]},{"slug":"flashinfer-ai-flashinfer","name":"flashinfer","tagline":"FlashInfer is a kernel library for serving large language models","github_url":"https://github.com/flashinfer-ai/flashinfer","owner":"flashinfer-ai","repo":"flashinfer","owner_avatar_url":"https://avatars.githubusercontent.com/u/145061914?v=4","primary_language":"Python","stars":6231,"forks":1327,"topics":["attention","cuda","distributed-inference","gpu","jit","large-large-models","llm-inference","moe","nvidia","pytorch"],"archived":false,"github_pushed_at":"2026-08-24T17:00:11+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/flashinfer-ai-flashinfer","markdown_url":"https://www.graphcanon.com/tools/flashinfer-ai-flashinfer.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/flashinfer-ai-flashinfer","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=flashinfer-ai-flashinfer","shared_categories":["inference-serving"]},{"slug":"pykeio-ort","name":"ort","tagline":"Fast ML inference and training for ONNX models in Rust","github_url":"https://github.com/pykeio/ort","owner":"pykeio","repo":"ort","owner_avatar_url":"https://avatars.githubusercontent.com/u/62268720?v=4","primary_language":"Rust","stars":2472,"forks":263,"topics":["ai","ai-training","fine-tuning","inference","machine-learning","onnx","onnxruntime","rust"],"archived":false,"github_pushed_at":"2026-08-23T20:46:26+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/pykeio-ort","markdown_url":"https://www.graphcanon.com/tools/pykeio-ort.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/pykeio-ort","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=pykeio-ort","shared_categories":["model-training","inference-serving"]},{"slug":"huangowen-awesome-llm-compression","name":"Awesome-LLM-Compression","tagline":"Awesome LLM compression research papers and tools to accelerate LLM training and inference.","github_url":"https://github.com/HuangOwen/Awesome-LLM-Compression","owner":"HuangOwen","repo":"Awesome-LLM-Compression","owner_avatar_url":"https://avatars.githubusercontent.com/u/24937399?v=4","primary_language":null,"stars":1859,"forks":129,"topics":[],"archived":false,"github_pushed_at":"2026-06-30T15:26:46+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/huangowen-awesome-llm-compression","markdown_url":"https://www.graphcanon.com/tools/huangowen-awesome-llm-compression.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/huangowen-awesome-llm-compression","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=huangowen-awesome-llm-compression","shared_categories":["inference-serving"]},{"slug":"tensorflow-mesh","name":"mesh","tagline":"Mesh TensorFlow: Model Parallelism Made Easier","github_url":"https://github.com/tensorflow/mesh","owner":"tensorflow","repo":"mesh","owner_avatar_url":"https://avatars.githubusercontent.com/u/15658638?v=4","primary_language":"Python","stars":1630,"forks":255,"topics":[],"archived":true,"github_pushed_at":"2023-11-17T19:39:54+00:00","maintenance_label":"Archived","url":"https://www.graphcanon.com/tools/tensorflow-mesh","markdown_url":"https://www.graphcanon.com/tools/tensorflow-mesh.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/tensorflow-mesh","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=tensorflow-mesh","shared_categories":["model-training"]},{"slug":"tensorflow-model-optimization","name":"model-optimization","tagline":"Toolkit for optimizing ML models in Keras and TensorFlow","github_url":"https://github.com/tensorflow/model-optimization","owner":"tensorflow","repo":"model-optimization","owner_avatar_url":"https://avatars.githubusercontent.com/u/15658638?v=4","primary_language":"Python","stars":1576,"forks":346,"topics":["compression","deep-learning","keras","machine-learning","ml","model-compression","optimization","pruning","quantization","quantized-networks","quantized-neural-networks","quantized-training","sparsity","tensorflow"],"archived":false,"github_pushed_at":"2026-07-27T06:03:18+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/tensorflow-model-optimization","markdown_url":"https://www.graphcanon.com/tools/tensorflow-model-optimization.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/tensorflow-model-optimization","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=tensorflow-model-optimization","shared_categories":["model-training"]},{"slug":"jmaczan-tiny-vllm","name":"tiny-vllm","tagline":"Build your own high performance LLM inference engine in C++ and CUDA - a smaller version of vLLM","github_url":"https://github.com/jmaczan/tiny-vllm","owner":"jmaczan","repo":"tiny-vllm","owner_avatar_url":"https://avatars.githubusercontent.com/u/18054202?v=4","primary_language":"C++","stars":1075,"forks":84,"topics":["ai","attention","batching","course","cpp","cuda","hpc","inference","llm","llm-inference","pagedattention","tiny-vllm","vllm"],"archived":false,"github_pushed_at":"2026-08-23T14:20:13+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/jmaczan-tiny-vllm","markdown_url":"https://www.graphcanon.com/tools/jmaczan-tiny-vllm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/jmaczan-tiny-vllm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=jmaczan-tiny-vllm","shared_categories":["inference-serving"]},{"slug":"andrewkchan-yalm","name":"yalm","tagline":"LLM inference engine in C++/CUDA without dependency on external libraries except for I/O","github_url":"https://github.com/andrewkchan/yalm","owner":"andrewkchan","repo":"yalm","owner_avatar_url":"https://avatars.githubusercontent.com/u/8591901?v=4","primary_language":"C++","stars":596,"forks":64,"topics":["cpp","cuda","inference-engine","llama","llamacpp","llm","llm-inference","machine-learning","mistral"],"archived":false,"github_pushed_at":"2025-09-13T09:22:40+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/andrewkchan-yalm","markdown_url":"https://www.graphcanon.com/tools/andrewkchan-yalm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/andrewkchan-yalm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=andrewkchan-yalm","shared_categories":["inference-serving"]}]}}