{"data":{"slug":"lyogavin-airllm","name":"airllm","tagline":"AirLLM 70B inference with single 4GB GPU","github_url":"https://github.com/lyogavin/airllm","owner":"lyogavin","repo":"airllm","owner_avatar_url":"https://avatars.githubusercontent.com/u/1113905?v=4","primary_language":"Jupyter Notebook","stars":24183,"forks":2722,"topics":["chinese-llm","chinese-nlp","finetune","generative-ai","instruct-gpt","instruction-set","llama","llm","lora","open-models","open-source","open-source-models","qlora"],"archived":false,"github_pushed_at":"2026-07-23T08:29:43+00:00","maintenance_label":"Very active","url":"https://www.graphcanon.com/tools/lyogavin-airllm","markdown_url":"https://www.graphcanon.com/tools/lyogavin-airllm.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/lyogavin-airllm","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=lyogavin-airllm","description":"AirLLM 70B inference with single 4GB GPU","homepage_url":null,"license":"Apache-2.0","open_issues":115,"watchers":235,"ai_summary":"A framework for running large language model (LLM) inferences on low-resource hardware, specifically a 4GB GPU.","readme_excerpt":"### 1. Install package\n\nFirst, install the airllm pip package.\n\n```bash\npip install airllm\n```","github_created_at":"2023-06-12T16:28:41+00:00","created_at":"2026-07-07T17:33:04.588292+00:00","updated_at":"2026-07-28T18:00:23.007156+00:00","categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"}],"tags":[{"slug":"chinese-llm","name":"chinese-llm"},{"slug":"chinese-nlp","name":"chinese-nlp"},{"slug":"finetune","name":"finetune"},{"slug":"generative-ai","name":"generative-ai"},{"slug":"instruct-gpt","name":"instruct-gpt"},{"slug":"instruction-set","name":"instruction-set"},{"slug":"llama","name":"llama"},{"slug":"llm","name":"llm"}],"trust":{"provenance":{"is_fork":false,"github_id":652712035,"owner_type":"User","methodology":"github_public_v1","parent_repo":null,"near_duplicate_slugs":[]},"computed_at":"2026-07-28T18:00:22.092Z","maintenance":{"label":"Very active","score":96,"methodology":"github_public_v1","releases_90d":2,"days_since_push":5,"last_release_at":"2026-06-30T23:02:38Z"},"security_summary":{"status":"findings","scanner":"osv@v1","low_count":4,"high_count":0,"last_scan_at":"2026-07-09T11:30:25.655Z","medium_count":0,"scan_profile":"deps","critical_count":0}},"capability_facts":{"scan":{"source":"repo_scan","observed_at":"2026-07-28T18:00:22.709Z"},"languages":{"value":["jupyter notebook"],"source":"github.language","observed_at":"2026-07-28T18:00:22.709Z"},"license_spdx":{"value":"Apache-2.0","source":"github.license","observed_at":"2026-07-28T18:00:22.709Z"}},"decision_facts":{"hosting":null,"pricing":{"model":"freemium","summary":"Free and open-source under the Apache-2.0 license; however, infrastructure costs apply."},"requirements":{"notes":["A single 4GB GPU is sufficient for using this framework to run large language model inferences."],"min_ram_gb":16,"requires_docker":false},"constraints":{"min_ram_gb":16,"pricing_model":"freemium","requires_docker":false},"when_to_use":["If you have limited hardware resources but need to perform inferences on large language models (like the 70B parameter model that AirLLM supports), use AirLLM.","AirLLM is ideal if your project involves Chinese NLP or LLMs, as it stands out with support for models like `chinese-llm`."],"when_not_to_use":["Avoid using AirLLM if you require models to run on higher-end GPUs or multiple GPU clusters, as its strength lies in low-resource efficiency.","Do not use AirLLM if you are working primarily with non-Chinese language datasets and models, since support for other languages may be less optimized compared to competition."],"source":"enrich:decision_facts","observed_at":"2026-07-11T00:49:02.990Z"},"constraint_facets":{"min_ram_gb":16,"pricing_model":"freemium","requires_docker":false},"decision_summary":[{"label":"Pricing","value":"freemium - Free and open-source under the Apache-2.0 license; however, infrastructure costs apply."},{"label":"Requirements","value":"Min 16 GB RAM; A single 4GB GPU is sufficient for using this framework to run large language model inferences."},{"label":"Adopt for","value":"AirLLM is a notable framework designed specifically for running large language models on low-resource hardware, such as a single 4GB GPU."},{"label":"License detail","value":"Apache-2.0"}]}}