{"data":{"node":{"slug":"huggingface-trl","name":"trl","tagline":"Train transformer language models with reinforcement learning.","github_url":"https://github.com/huggingface/trl","owner":"huggingface","repo":"trl","owner_avatar_url":"https://avatars.githubusercontent.com/u/25720743?v=4","primary_language":"Python","stars":19016,"forks":2891,"topics":[],"archived":false,"github_pushed_at":"2026-08-06T10:02:43+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/huggingface-trl","markdown_url":"https://www.graphcanon.com/tools/huggingface-trl.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/huggingface-trl","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=huggingface-trl"},"categories":[{"slug":"model-training","name":"Model Training","url":"https://www.graphcanon.com/categories/model-training","markdown_url":"https://www.graphcanon.com/categories/model-training.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/model-training"}],"tags":[{"slug":"distributed-training","name":"distributed-training"},{"slug":"reinforcement-learning","name":"reinforcement-learning"},{"slug":"transformers","name":"transformers"}],"edges":[{"type":"related","direction":"out","explanation":"TRL is a specific library for training transformer models with reinforcement learning, while awesome-RLHF is a curated list of resources on RLHF. They are related because TRL provides practical tools in the space covered by the resource list.","successor_context":null,"tool":{"slug":"opendilab-awesome-rlhf","name":"awesome-RLHF","tagline":"A curated list of reinforcement learning with human feedback resources (continually updated)","github_url":"https://github.com/opendilab/awesome-RLHF","owner":"opendilab","repo":"awesome-RLHF","owner_avatar_url":"https://avatars.githubusercontent.com/u/86840398?v=4","primary_language":null,"stars":4422,"forks":258,"topics":["deep-learning","deep-reinforcement-learning","human-feedback","large-language-models","reinforcement-learning","rlhf"],"archived":false,"github_pushed_at":"2026-05-20T12:56:15+00:00","maintenance_label":"Steady","stars_delta_30d":9,"url":"https://www.graphcanon.com/tools/opendilab-awesome-rlhf","markdown_url":"https://www.graphcanon.com/tools/opendilab-awesome-rlhf.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/opendilab-awesome-rlhf","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=opendilab-awesome-rlhf"}},{"type":"integrates_with","direction":"out","explanation":"TRL integrates with transformers by leveraging transformers' extensive library of pre-trained models as a foundational base for applying reinforcement learning methods such as SFT, GRPO, and DPO to further refine these models. This relationship allows TRL to utilize the robust model architectures provided by transformers while focusing on advanced training techniques.","successor_context":null,"tool":{"slug":"huggingface-transformers","name":"transformers","tagline":"Transformers: the model-definition framework for state-of-the-art machine learning models in text, vision, audio, and multimodal models","github_url":"https://github.com/huggingface/transformers","owner":"huggingface","repo":"transformers","owner_avatar_url":"https://avatars.githubusercontent.com/u/25720743?v=4","primary_language":"Python","stars":164121,"forks":34249,"topics":["audio","deep-learning","deepseek","gemma","glm","hacktoberfest","llm","machine-learning","model-hub","natural-language-processing","nlp","pretrained-models","python","pytorch","pytorch-transformers","qwen","speech-recognition","transformer","vlm"],"archived":false,"github_pushed_at":"2026-08-15T22:28:12+00:00","maintenance_label":"Very active","stars_delta_30d":1457,"url":"https://www.graphcanon.com/tools/huggingface-transformers","markdown_url":"https://www.graphcanon.com/tools/huggingface-transformers.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/huggingface-transformers","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=huggingface-transformers"}}],"neighbours":[{"slug":"huggingface-peft","name":"peft","tagline":"State-of-the-art Parameter-Efficient Fine-Tuning","github_url":"https://github.com/huggingface/peft","owner":"huggingface","repo":"peft","owner_avatar_url":"https://avatars.githubusercontent.com/u/25720743?v=4","primary_language":"Python","stars":21443,"forks":2399,"topics":["adapter","diffusion","fine-tuning","llm","lora","parameter-efficient-learning","peft","python","pytorch","transformers"],"archived":false,"github_pushed_at":"2026-07-23T15:52:27+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/huggingface-peft","markdown_url":"https://www.graphcanon.com/tools/huggingface-peft.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/huggingface-peft","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=huggingface-peft","shared_categories":["model-training"]},{"slug":"nvidia-megatron-lm","name":"Megatron-LM","tagline":"Ongoing research training transformer models at scale","github_url":"https://github.com/NVIDIA/Megatron-LM","owner":"NVIDIA","repo":"Megatron-LM","owner_avatar_url":"https://avatars.githubusercontent.com/u/1728152?v=4","primary_language":"Python","stars":17341,"forks":4333,"topics":["large-language-models","model-para","transformers"],"archived":false,"github_pushed_at":"2026-08-06T23:12:52+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/nvidia-megatron-lm","markdown_url":"https://www.graphcanon.com/tools/nvidia-megatron-lm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/nvidia-megatron-lm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=nvidia-megatron-lm","shared_categories":["model-training"]},{"slug":"nvidia-tensorrt-llm","name":"TensorRT-LLM","tagline":"Python API for defining and optimizing Large Language Models (LLMs) on NVIDIA GPUs","github_url":"https://github.com/NVIDIA/TensorRT-LLM","owner":"NVIDIA","repo":"TensorRT-LLM","owner_avatar_url":"https://avatars.githubusercontent.com/u/1728152?v=4","primary_language":"Python","stars":14317,"forks":2641,"topics":["blackwell","cuda","llm-serving","moe","pytorch"],"archived":false,"github_pushed_at":"2026-08-07T05:40:26+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/nvidia-tensorrt-llm","markdown_url":"https://www.graphcanon.com/tools/nvidia-tensorrt-llm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/nvidia-tensorrt-llm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=nvidia-tensorrt-llm","shared_categories":[]},{"slug":"lightning-ai-litgpt","name":"litgpt","tagline":"High-performance LLMs with recipes for pretraining, finetuning and deployment","github_url":"https://github.com/Lightning-AI/litgpt","owner":"Lightning-AI","repo":"litgpt","owner_avatar_url":"https://avatars.githubusercontent.com/u/58386951?v=4","primary_language":"Python","stars":13605,"forks":1483,"topics":["ai","artificial-intelligence","deep-learning","large-language-models","llm","llm-inference","llms"],"archived":false,"github_pushed_at":"2026-07-20T10:24:12+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/lightning-ai-litgpt","markdown_url":"https://www.graphcanon.com/tools/lightning-ai-litgpt.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/lightning-ai-litgpt","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=lightning-ai-litgpt","shared_categories":["model-training"]},{"slug":"shishirpatil-gorilla","name":"gorilla","tagline":"Training and Evaluating LLMs for Function Calls (Tool Calls)","github_url":"https://github.com/ShishirPatil/gorilla","owner":"ShishirPatil","repo":"gorilla","owner_avatar_url":"https://avatars.githubusercontent.com/u/30296397?v=4","primary_language":"Python","stars":12988,"forks":1397,"topics":["api","api-documentation","chatgpt","claude-api","gpt-4-api","llm","openai-api","openai-functions"],"archived":false,"github_pushed_at":"2026-04-13T03:19:45+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/shishirpatil-gorilla","markdown_url":"https://www.graphcanon.com/tools/shishirpatil-gorilla.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/shishirpatil-gorilla","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=shishirpatil-gorilla","shared_categories":["model-training"]},{"slug":"fareedkhan-dev-train-llm-from-scratch","name":"train-llm-from-scratch","tagline":"A straightforward method for training your LLM from raw text to aligned model generation","github_url":"https://github.com/FareedKhan-dev/train-llm-from-scratch","owner":"FareedKhan-dev","repo":"train-llm-from-scratch","owner_avatar_url":"https://avatars.githubusercontent.com/u/63067900?v=4","primary_language":"Python","stars":9141,"forks":1264,"topics":["gemini","large-language-models","llm","openai","training","transformers"],"archived":false,"github_pushed_at":"2026-08-17T05:07:26+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/fareedkhan-dev-train-llm-from-scratch","markdown_url":"https://www.graphcanon.com/tools/fareedkhan-dev-train-llm-from-scratch.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/fareedkhan-dev-train-llm-from-scratch","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=fareedkhan-dev-train-llm-from-scratch","shared_categories":["model-training"]},{"slug":"nvidia-fastertransformer","name":"FasterTransformer","tagline":"Transformer related optimization including BERT and GPT","github_url":"https://github.com/NVIDIA/FasterTransformer","owner":"NVIDIA","repo":"FasterTransformer","owner_avatar_url":"https://avatars.githubusercontent.com/u/1728152?v=4","primary_language":"C++","stars":6446,"forks":935,"topics":["bert","gpt","pytorch","transformer"],"archived":false,"github_pushed_at":"2024-03-27T11:25:30+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/nvidia-fastertransformer","markdown_url":"https://www.graphcanon.com/tools/nvidia-fastertransformer.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/nvidia-fastertransformer","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=nvidia-fastertransformer","shared_categories":[]},{"slug":"meta-pytorch-torchtune","name":"torchtune","tagline":"PyTorch native post-training library","github_url":"https://github.com/meta-pytorch/torchtune","owner":"meta-pytorch","repo":"torchtune","owner_avatar_url":"https://avatars.githubusercontent.com/u/107212512?v=4","primary_language":"Python","stars":5793,"forks":743,"topics":[],"archived":false,"github_pushed_at":"2026-08-06T12:15:22+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/meta-pytorch-torchtune","markdown_url":"https://www.graphcanon.com/tools/meta-pytorch-torchtune.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/meta-pytorch-torchtune","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=meta-pytorch-torchtune","shared_categories":["model-training"]},{"slug":"opendilab-awesome-rlhf","name":"awesome-RLHF","tagline":"A curated list of reinforcement learning with human feedback resources (continually updated)","github_url":"https://github.com/opendilab/awesome-RLHF","owner":"opendilab","repo":"awesome-RLHF","owner_avatar_url":"https://avatars.githubusercontent.com/u/86840398?v=4","primary_language":null,"stars":4422,"forks":258,"topics":["deep-learning","deep-reinforcement-learning","human-feedback","large-language-models","reinforcement-learning","rlhf"],"archived":false,"github_pushed_at":"2026-05-20T12:56:15+00:00","maintenance_label":"Slowing","url":"https://www.graphcanon.com/tools/opendilab-awesome-rlhf","markdown_url":"https://www.graphcanon.com/tools/opendilab-awesome-rlhf.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/opendilab-awesome-rlhf","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=opendilab-awesome-rlhf","shared_categories":["model-training"]},{"slug":"truera-trulens","name":"trulens","tagline":"Evaluation and Tracking for LLM Experiments and AI Agents","github_url":"https://github.com/truera/trulens","owner":"truera","repo":"trulens","owner_avatar_url":"https://avatars.githubusercontent.com/u/51224128?v=4","primary_language":"Python","stars":3516,"forks":327,"topics":["agent-evaluation","agentops","ai-agents","ai-monitoring","ai-observability","evals","explainable-ml","llm-eval","llm-evaluation","llmops","llms","machine-learning","neural-networks"],"archived":false,"github_pushed_at":"2026-08-20T10:21:00+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/truera-trulens","markdown_url":"https://www.graphcanon.com/tools/truera-trulens.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/truera-trulens","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=truera-trulens","shared_categories":[]},{"slug":"alibaba-roll","name":"ROLL","tagline":"Scaling Library for Reinforcement Learning with Large Language Models","github_url":"https://github.com/alibaba/ROLL","owner":"alibaba","repo":"ROLL","owner_avatar_url":"https://avatars.githubusercontent.com/u/1961952?v=4","primary_language":"Python","stars":3354,"forks":304,"topics":["agentic","rlhf","rlvr"],"archived":false,"github_pushed_at":"2026-08-07T01:51:58+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/alibaba-roll","markdown_url":"https://www.graphcanon.com/tools/alibaba-roll.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/alibaba-roll","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=alibaba-roll","shared_categories":["model-training"]},{"slug":"ashishpatel26-llm-finetuning","name":"LLM-Finetuning","tagline":"LLM Finetuning with PEFT","github_url":"https://github.com/ashishpatel26/LLM-Finetuning","owner":"ashishpatel26","repo":"LLM-Finetuning","owner_avatar_url":"https://avatars.githubusercontent.com/u/3095771?v=4","primary_language":"Jupyter Notebook","stars":2966,"forks":769,"topics":["falcon","fine-tuning","huggingface","llama","llama2","llm","llms","lora","peft","pytorch","text-generation"],"archived":false,"github_pushed_at":"2025-08-01T12:00:20+00:00","maintenance_label":"Dormant","url":"https://www.graphcanon.com/tools/ashishpatel26-llm-finetuning","markdown_url":"https://www.graphcanon.com/tools/ashishpatel26-llm-finetuning.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ashishpatel26-llm-finetuning","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ashishpatel26-llm-finetuning","shared_categories":["model-training"]}]}}