{"data":{"node":{"slug":"optimalscale-lmflow","name":"LMFlow","tagline":"An Extensible Toolkit for Finetuning and Inference of Large Foundation Models","github_url":"https://github.com/OptimalScale/LMFlow","owner":"OptimalScale","repo":"LMFlow","owner_avatar_url":"https://avatars.githubusercontent.com/u/128913633?v=4","primary_language":"Python","stars":8486,"forks":825,"topics":["chatgpt","deep-learning","instruction-following","language-model","pretrained-models","pytorch","transformer"],"archived":false,"github_pushed_at":"2026-05-22T02:57:26+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/optimalscale-lmflow","markdown_url":"https://www.graphcanon.com/tools/optimalscale-lmflow.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/optimalscale-lmflow","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=optimalscale-lmflow"},"categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"},{"slug":"llm-frameworks","name":"LLM Frameworks","url":"https://www.graphcanon.com/categories/llm-frameworks","markdown_url":"https://www.graphcanon.com/categories/llm-frameworks.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/llm-frameworks"}],"tags":[{"slug":"chatgpt","name":"chatgpt"},{"slug":"deep-learning","name":"deep-learning"},{"slug":"instruction-following","name":"instruction-following"},{"slug":"language-model","name":"language-model"},{"slug":"pretrained-models","name":"pretrained-models"},{"slug":"pytorch","name":"pytorch"},{"slug":"transformer","name":"transformer"}],"edges":[],"neighbours":[{"slug":"lyogavin-airllm","name":"airllm","tagline":"AirLLM 70B inference with single 4GB GPU","github_url":"https://github.com/lyogavin/airllm","owner":"lyogavin","repo":"airllm","owner_avatar_url":"https://avatars.githubusercontent.com/u/1113905?v=4","primary_language":"Jupyter Notebook","stars":24183,"forks":2722,"topics":["chinese-llm","chinese-nlp","finetune","generative-ai","instruct-gpt","instruction-set","llama","llm","lora","open-models","open-source","open-source-models","qlora"],"archived":false,"github_pushed_at":"2026-07-23T08:29:43+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/lyogavin-airllm","markdown_url":"https://www.graphcanon.com/tools/lyogavin-airllm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/lyogavin-airllm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=lyogavin-airllm","shared_categories":["inference-serving"]},{"slug":"huggingface-peft","name":"peft","tagline":"State-of-the-art Parameter-Efficient Fine-Tuning","github_url":"https://github.com/huggingface/peft","owner":"huggingface","repo":"peft","owner_avatar_url":"https://avatars.githubusercontent.com/u/25720743?v=4","primary_language":"Python","stars":21585,"forks":2446,"topics":["adapter","diffusion","fine-tuning","llm","lora","parameter-efficient-learning","peft","python","pytorch","transformers"],"archived":false,"github_pushed_at":"2026-08-22T01:15:06+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/huggingface-peft","markdown_url":"https://www.graphcanon.com/tools/huggingface-peft.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/huggingface-peft","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=huggingface-peft","shared_categories":["llm-frameworks"]},{"slug":"lightning-ai-litgpt","name":"litgpt","tagline":"High-performance LLMs with recipes for pretraining, finetuning and deployment","github_url":"https://github.com/Lightning-AI/litgpt","owner":"Lightning-AI","repo":"litgpt","owner_avatar_url":"https://avatars.githubusercontent.com/u/58386951?v=4","primary_language":"Python","stars":13605,"forks":1483,"topics":["ai","artificial-intelligence","deep-learning","large-language-models","llm","llm-inference","llms"],"archived":false,"github_pushed_at":"2026-07-20T10:24:12+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/lightning-ai-litgpt","markdown_url":"https://www.graphcanon.com/tools/lightning-ai-litgpt.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/lightning-ai-litgpt","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=lightning-ai-litgpt","shared_categories":["llm-frameworks","inference-serving"]},{"slug":"shishirpatil-gorilla","name":"gorilla","tagline":"Training and Evaluating LLMs for Function Calls (Tool Calls)","github_url":"https://github.com/ShishirPatil/gorilla","owner":"ShishirPatil","repo":"gorilla","owner_avatar_url":"https://avatars.githubusercontent.com/u/30296397?v=4","primary_language":"Python","stars":12988,"forks":1397,"topics":["api","api-documentation","chatgpt","claude-api","gpt-4-api","llm","openai-api","openai-functions"],"archived":false,"github_pushed_at":"2026-04-13T03:19:45+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/shishirpatil-gorilla","markdown_url":"https://www.graphcanon.com/tools/shishirpatil-gorilla.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/shishirpatil-gorilla","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=shishirpatil-gorilla","shared_categories":[]},{"slug":"fminference-flexllmgen","name":"FlexLLMGen","tagline":"Running large language models on a single GPU for throughput-oriented scenarios.","github_url":"https://github.com/FMInference/FlexLLMGen","owner":"FMInference","repo":"FlexLLMGen","owner_avatar_url":"https://avatars.githubusercontent.com/u/125944572?v=4","primary_language":"Python","stars":9361,"forks":590,"topics":["deep-learning","gpt-3","high-throughput","large-language-models","machine-learning","offloading","opt"],"archived":true,"github_pushed_at":"2024-10-28T03:05:41+00:00","maintenance_label":"Archived","url":"https://www.graphcanon.com/tools/fminference-flexllmgen","markdown_url":"https://www.graphcanon.com/tools/fminference-flexllmgen.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/fminference-flexllmgen","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=fminference-flexllmgen","shared_categories":["inference-serving"]},{"slug":"fareedkhan-dev-train-llm-from-scratch","name":"train-llm-from-scratch","tagline":"A straightforward method for training your LLM from raw text to aligned model generation","github_url":"https://github.com/FareedKhan-dev/train-llm-from-scratch","owner":"FareedKhan-dev","repo":"train-llm-from-scratch","owner_avatar_url":"https://avatars.githubusercontent.com/u/63067900?v=4","primary_language":"Python","stars":9141,"forks":1264,"topics":["gemini","large-language-models","llm","openai","training","transformers"],"archived":false,"github_pushed_at":"2026-08-17T05:07:26+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/fareedkhan-dev-train-llm-from-scratch","markdown_url":"https://www.graphcanon.com/tools/fareedkhan-dev-train-llm-from-scratch.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/fareedkhan-dev-train-llm-from-scratch","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=fareedkhan-dev-train-llm-from-scratch","shared_categories":["inference-serving"]},{"slug":"flashinfer-ai-flashinfer","name":"flashinfer","tagline":"FlashInfer is a kernel library for serving large language models","github_url":"https://github.com/flashinfer-ai/flashinfer","owner":"flashinfer-ai","repo":"flashinfer","owner_avatar_url":"https://avatars.githubusercontent.com/u/145061914?v=4","primary_language":"Python","stars":6231,"forks":1327,"topics":["attention","cuda","distributed-inference","gpu","jit","large-large-models","llm-inference","moe","nvidia","pytorch"],"archived":false,"github_pushed_at":"2026-08-24T17:00:11+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/flashinfer-ai-flashinfer","markdown_url":"https://www.graphcanon.com/tools/flashinfer-ai-flashinfer.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/flashinfer-ai-flashinfer","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=flashinfer-ai-flashinfer","shared_categories":["llm-frameworks","inference-serving"]},{"slug":"ashishpatel26-llm-finetuning","name":"LLM-Finetuning","tagline":"LLM Finetuning with PEFT","github_url":"https://github.com/ashishpatel26/LLM-Finetuning","owner":"ashishpatel26","repo":"LLM-Finetuning","owner_avatar_url":"https://avatars.githubusercontent.com/u/3095771?v=4","primary_language":"Jupyter Notebook","stars":2979,"forks":771,"topics":["falcon","fine-tuning","huggingface","llama","llama2","llm","llms","lora","peft","pytorch","text-generation"],"archived":false,"github_pushed_at":"2025-08-01T12:00:20+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/ashishpatel26-llm-finetuning","markdown_url":"https://www.graphcanon.com/tools/ashishpatel26-llm-finetuning.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ashishpatel26-llm-finetuning","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ashishpatel26-llm-finetuning","shared_categories":["llm-frameworks"]},{"slug":"stanford-crfm-helm","name":"helm","tagline":"Holistic, reproducible and transparent evaluation of foundation models","github_url":"https://github.com/stanford-crfm/helm","owner":"stanford-crfm","repo":"helm","owner_avatar_url":"https://avatars.githubusercontent.com/u/75054807?v=4","primary_language":"Python","stars":2873,"forks":406,"topics":[],"archived":false,"github_pushed_at":"2026-08-01T01:23:17+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/stanford-crfm-helm","markdown_url":"https://www.graphcanon.com/tools/stanford-crfm-helm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/stanford-crfm-helm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=stanford-crfm-helm","shared_categories":[]},{"slug":"kubeflow-trainer","name":"trainer","tagline":"Distributed AI Model Training and LLM Fine-Tuning on Kubernetes","github_url":"https://github.com/kubeflow/trainer","owner":"kubeflow","repo":"trainer","owner_avatar_url":"https://avatars.githubusercontent.com/u/33164907?v=4","primary_language":"Go","stars":2196,"forks":1030,"topics":["ai","distributed","fine-tuning","gpu","huggingface","jax","kubeflow","kubernetes","llm","machine-learning","mlops","python","pytorch","tensorflow","xgboost"],"archived":false,"github_pushed_at":"2026-08-22T02:27:28+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/kubeflow-trainer","markdown_url":"https://www.graphcanon.com/tools/kubeflow-trainer.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/kubeflow-trainer","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=kubeflow-trainer","shared_categories":["llm-frameworks"]},{"slug":"xllm-ai-xllm","name":"xllm","tagline":"A high-performance inference engine for LLM, VLM, DiT and REC models","github_url":"https://github.com/xLLM-AI/xllm","owner":"xLLM-AI","repo":"xllm","owner_avatar_url":"https://avatars.githubusercontent.com/u/205719415?v=4","primary_language":"C++","stars":1534,"forks":282,"topics":["deepseek","glm","inference","inference-engine","large-language-models","llm-inference","qwen"],"archived":false,"github_pushed_at":"2026-08-24T09:51:10+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/xllm-ai-xllm","markdown_url":"https://www.graphcanon.com/tools/xllm-ai-xllm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/xllm-ai-xllm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=xllm-ai-xllm","shared_categories":["inference-serving"]},{"slug":"arahim3-mlx-tune","name":"mlx-tune","tagline":"Fine-tune LLMs on your Mac with Apple Silicon for various tasks including SFT, DPO, GRPO, Vision, TTS, STT, Embedding, and OCR.","github_url":"https://github.com/ARahim3/mlx-tune","owner":"ARahim3","repo":"mlx-tune","owner_avatar_url":"https://avatars.githubusercontent.com/u/41390319?v=4","primary_language":"Python","stars":1372,"forks":88,"topics":["apple-silicon","deep-learning","huggingface","large-language-models","llm","llm-finetuning","local-llm","lora","machine-learning","macos","mlx","on-device-ai","peft","speech-recognition","speech-to-text","text-to-speech","transformers","unsloth","vision-language-model","whisper"],"archived":false,"github_pushed_at":"2026-06-23T12:24:30+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/arahim3-mlx-tune","markdown_url":"https://www.graphcanon.com/tools/arahim3-mlx-tune.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/arahim3-mlx-tune","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=arahim3-mlx-tune","shared_categories":["llm-frameworks"]}]}}