{"data":{"node":{"slug":"ray-project-ray-llm","name":"ray-llm","tagline":"Archived repository; LLM serving APIs integrated into the Ray project","github_url":"https://github.com/ray-project/ray-llm","owner":"ray-project","repo":"ray-llm","owner_avatar_url":"https://avatars.githubusercontent.com/u/22125274?v=4","primary_language":null,"stars":1261,"forks":90,"topics":["llm","llm-serving","ray"],"archived":true,"github_pushed_at":"2025-03-13T01:13:38+00:00","maintenance_label":"Archived","url":"https://www.graphcanon.com/tools/ray-project-ray-llm","markdown_url":"https://www.graphcanon.com/tools/ray-project-ray-llm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ray-project-ray-llm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ray-project-ray-llm"},"categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"},{"slug":"model-training","name":"Model Training","url":"https://www.graphcanon.com/categories/model-training","markdown_url":"https://www.graphcanon.com/categories/model-training.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/model-training"}],"tags":[{"slug":"llm-serving","name":"llm-serving"},{"slug":"ray","name":"ray"}],"edges":[],"neighbours":[{"slug":"vllm-project-vllm","name":"vllm","tagline":"A high-throughput and memory-efficient inference and serving engine for LLMs","github_url":"https://github.com/vllm-project/vllm","owner":"vllm-project","repo":"vllm","owner_avatar_url":"https://avatars.githubusercontent.com/u/136984999?v=4","primary_language":"Python","stars":87847,"forks":20135,"topics":["amd","blackwell","cuda","deepseek","deepseek-v3","gpt","gpt-oss","inference","kimi","llama","llm","llm-serving","model-serving","moe","openai","pytorch","qwen","qwen3","tpu","transformer"],"archived":false,"github_pushed_at":"2026-08-01T11:55:36+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/vllm-project-vllm","markdown_url":"https://www.graphcanon.com/tools/vllm-project-vllm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/vllm-project-vllm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=vllm-project-vllm","shared_categories":["inference-serving"]},{"slug":"berriai-litellm","name":"litellm","tagline":"Python SDK and Proxy Server for calling multiple LLM APIs","github_url":"https://github.com/BerriAI/litellm","owner":"BerriAI","repo":"litellm","owner_avatar_url":"https://avatars.githubusercontent.com/u/121462774?v=4","primary_language":"Python","stars":55221,"forks":10231,"topics":["ai-gateway","anthropic","azure-openai","bedrock","gateway","langchain","litellm","llm","llm-gateway","llmops","mcp-gateway","openai","openai-proxy","rust","rust-ai","vertex-ai"],"archived":false,"github_pushed_at":"2026-08-01T05:53:28+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/berriai-litellm","markdown_url":"https://www.graphcanon.com/tools/berriai-litellm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/berriai-litellm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=berriai-litellm","shared_categories":["inference-serving"]},{"slug":"lightning-ai-litgpt","name":"litgpt","tagline":"High-performance LLMs with recipes for pretraining, finetuning and deployment","github_url":"https://github.com/Lightning-AI/litgpt","owner":"Lightning-AI","repo":"litgpt","owner_avatar_url":"https://avatars.githubusercontent.com/u/58386951?v=4","primary_language":"Python","stars":13605,"forks":1483,"topics":["ai","artificial-intelligence","deep-learning","large-language-models","llm","llm-inference","llms"],"archived":false,"github_pushed_at":"2026-07-20T10:24:12+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/lightning-ai-litgpt","markdown_url":"https://www.graphcanon.com/tools/lightning-ai-litgpt.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/lightning-ai-litgpt","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=lightning-ai-litgpt","shared_categories":["model-training","inference-serving"]},{"slug":"wangrongsheng-awesome-llm-resources","name":"awesome-LLM-resources","tagline":"Summary of the world's best LLM resources.","github_url":"https://github.com/WangRongsheng/awesome-LLM-resources","owner":"WangRongsheng","repo":"awesome-LLM-resources","owner_avatar_url":"https://avatars.githubusercontent.com/u/55651568?v=4","primary_language":null,"stars":8845,"forks":950,"topics":["awesome-list","book","course","large-language-models","llama","llm","mistral","openai","qwen","rag","retrieval-augmented-generation","webui"],"archived":false,"github_pushed_at":"2026-08-14T15:54:28+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/wangrongsheng-awesome-llm-resources","markdown_url":"https://www.graphcanon.com/tools/wangrongsheng-awesome-llm-resources.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/wangrongsheng-awesome-llm-resources","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=wangrongsheng-awesome-llm-resources","shared_categories":["model-training","inference-serving"]},{"slug":"luhengshiwo-llmforeverybody","name":"LLMForEverybody","tagline":"LLM knowledge sharing for everyone, essential reading before big model interviews","github_url":"https://github.com/luhengshiwo/LLMForEverybody","owner":"luhengshiwo","repo":"LLMForEverybody","owner_avatar_url":"https://avatars.githubusercontent.com/u/13251733?v=4","primary_language":"Jupyter Notebook","stars":7167,"forks":662,"topics":["agent","interview-practice","interview-questions","learnllm","llm","rag"],"archived":false,"github_pushed_at":"2026-08-17T02:41:01+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/luhengshiwo-llmforeverybody","markdown_url":"https://www.graphcanon.com/tools/luhengshiwo-llmforeverybody.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/luhengshiwo-llmforeverybody","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=luhengshiwo-llmforeverybody","shared_categories":["model-training"]},{"slug":"andyyyy64-whichllm","name":"whichllm","tagline":"Command-line tool to find and benchmark local LLM performance","github_url":"https://github.com/Andyyyy64/whichllm","owner":"Andyyyy64","repo":"whichllm","owner_avatar_url":"https://avatars.githubusercontent.com/u/105579829?v=4","primary_language":"Python","stars":6225,"forks":330,"topics":["ai","apple-silicon","benchmarks","cli","command-line-tool","gguf","gpu","huggingface","inference","llm","local-llm","ollama","python","vram"],"archived":false,"github_pushed_at":"2026-08-05T07:15:32+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/andyyyy64-whichllm","markdown_url":"https://www.graphcanon.com/tools/andyyyy64-whichllm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/andyyyy64-whichllm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=andyyyy64-whichllm","shared_categories":["inference-serving"]},{"slug":"lazyagi-lazyllm","name":"LazyLLM","tagline":"Easiest and laziest way for building multi-agent LLMs applications.","github_url":"https://github.com/LazyAGI/LazyLLM","owner":"LazyAGI","repo":"LazyLLM","owner_avatar_url":"https://avatars.githubusercontent.com/u/171651681?v=4","primary_language":"Python","stars":3866,"forks":404,"topics":["agents","ai-agent","data","deep-learning","documentation-tool","finetuning","framework","knowlege-graph","langchain","lazyllm","llamaindex","llm","llms","rag"],"archived":false,"github_pushed_at":"2026-08-07T03:03:46+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/lazyagi-lazyllm","markdown_url":"https://www.graphcanon.com/tools/lazyagi-lazyllm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/lazyagi-lazyllm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=lazyagi-lazyllm","shared_categories":["model-training"]},{"slug":"b4rtaz-distributed-llama","name":"distributed-llama","tagline":"Distributed LLM inference using home devices cluster","github_url":"https://github.com/b4rtaz/distributed-llama","owner":"b4rtaz","repo":"distributed-llama","owner_avatar_url":"https://avatars.githubusercontent.com/u/12797776?v=4","primary_language":"C++","stars":3044,"forks":246,"topics":["distributed-computing","distributed-llm","llama2","llama3","llm","llm-inference","llms","neural-network","open-llm"],"archived":false,"github_pushed_at":"2026-07-05T16:47:20+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/b4rtaz-distributed-llama","markdown_url":"https://www.graphcanon.com/tools/b4rtaz-distributed-llama.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/b4rtaz-distributed-llama","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=b4rtaz-distributed-llama","shared_categories":["inference-serving"]},{"slug":"ray-project-llm-applications","name":"llm-applications","tagline":"Comprehensive guide to building RAG-based LLM applications for production","github_url":"https://github.com/ray-project/llm-applications","owner":"ray-project","repo":"llm-applications","owner_avatar_url":"https://avatars.githubusercontent.com/u/22125274?v=4","primary_language":"Jupyter Notebook","stars":1855,"forks":256,"topics":["anyscale","fine-tuning","llama2","llms","machine-learning","openai","ray","serving"],"archived":false,"github_pushed_at":"2026-08-15T00:14:04+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/ray-project-llm-applications","markdown_url":"https://www.graphcanon.com/tools/ray-project-llm-applications.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ray-project-llm-applications","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ray-project-llm-applications","shared_categories":["inference-serving"]},{"slug":"waybarrios-vllm-mlx","name":"vllm-mlx","tagline":"Server for LLMs and vision-language models compatible with Apple Silicon","github_url":"https://github.com/waybarrios/vllm-mlx","owner":"waybarrios","repo":"vllm-mlx","owner_avatar_url":"https://avatars.githubusercontent.com/u/6794828?v=4","primary_language":"Python","stars":1472,"forks":205,"topics":["anthropic","apple-silicon","audio-processing","claude-code","computer-vision","image-understanding","inference","llm","machine-learning","macos","mllm","mlx","multimodal-ai","speech-to-text","stt","text-to-speech","tts","video-understanding","vision-language-model","vllm"],"archived":false,"github_pushed_at":"2026-06-28T20:18:31+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/waybarrios-vllm-mlx","markdown_url":"https://www.graphcanon.com/tools/waybarrios-vllm-mlx.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/waybarrios-vllm-mlx","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=waybarrios-vllm-mlx","shared_categories":["model-training","inference-serving"]},{"slug":"benman1-generative-ai-with-langchain","name":"generative_ai_with_langchain","tagline":"Build production-ready LLM applications and advanced agents using Python, LangChain, and LangGraph","github_url":"https://github.com/benman1/generative_ai_with_langchain","owner":"benman1","repo":"generative_ai_with_langchain","owner_avatar_url":"https://avatars.githubusercontent.com/u/10786684?v=4","primary_language":"Jupyter Notebook","stars":1400,"forks":582,"topics":["agent","chatgpt","claude","claude-3-5-sonnet","deepseek","deepseek-r1","gpt","gpt-4o","huggingface","langchain","langgraph","llamacpp","llms","ollama","openai"],"archived":false,"github_pushed_at":"2026-08-05T12:50:30+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/benman1-generative-ai-with-langchain","markdown_url":"https://www.graphcanon.com/tools/benman1-generative-ai-with-langchain.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/benman1-generative-ai-with-langchain","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=benman1-generative-ai-with-langchain","shared_categories":[]},{"slug":"harleyszhang-llm-note","name":"llm_note","tagline":"LLM notes covering model inference transformer structures and framework analysis","github_url":"https://github.com/harleyszhang/llm_note","owner":"harleyszhang","repo":"llm_note","owner_avatar_url":"https://avatars.githubusercontent.com/u/37138671?v=4","primary_language":"Python","stars":888,"forks":90,"topics":["cuda-programming","kv-cache","llm","llm-inference","transformer-models","triton-kernels","vllm"],"archived":false,"github_pushed_at":"2026-08-19T06:46:41+00:00","maintenance_label":"Active","url":"https://www.graphcanon.com/tools/harleyszhang-llm-note","markdown_url":"https://www.graphcanon.com/tools/harleyszhang-llm-note.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/harleyszhang-llm-note","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=harleyszhang-llm-note","shared_categories":["inference-serving"]}]}}