{"data":{"slug":"b4rtaz-distributed-llama","name":"distributed-llama","tagline":"Distributed LLM inference using home devices cluster","github_url":"https://github.com/b4rtaz/distributed-llama","owner":"b4rtaz","repo":"distributed-llama","owner_avatar_url":"https://avatars.githubusercontent.com/u/12797776?v=4","primary_language":"C++","stars":3044,"forks":246,"topics":["distributed-computing","distributed-llm","llama2","llama3","llm","llm-inference","llms","neural-network","open-llm"],"archived":false,"github_pushed_at":"2026-07-05T16:47:20+00:00","maintenance_label":"Steady","stars_delta_30d":32,"url":"https://www.graphcanon.com/tools/b4rtaz-distributed-llama","markdown_url":"https://www.graphcanon.com/tools/b4rtaz-distributed-llama.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/b4rtaz-distributed-llama","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=b4rtaz-distributed-llama","description":"Distributed LLM inference. Connect home devices into a powerful cluster to accelerate LLM inference. More devices means faster inference.","homepage_url":null,"license":"MIT","open_issues":48,"watchers":52,"ai_summary":"A framework for distributing large language model inference across multiple connected home devices to boost performance","readme_excerpt":"## 💡 License\n\nThis project is released under the MIT license.","github_created_at":"2023-12-04T23:36:06+00:00","created_at":"2026-07-11T11:43:04.72866+00:00","updated_at":"2026-08-24T18:01:34.680249+00:00","categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"}],"tags":[{"slug":"distributed-computing","name":"distributed-computing"},{"slug":"llm-inference","name":"llm-inference"},{"slug":"neural-network","name":"neural-network"}],"trust":{"provenance":{"is_fork":false,"github_id":727470807,"owner_type":"User","methodology":"github_public_v1","parent_repo":null,"near_duplicate_slugs":[]},"computed_at":"2026-08-24T18:01:33.899Z","maintenance":{"label":"Steady","score":60,"methodology":"github_public_v1","releases_90d":0,"days_since_push":50,"last_release_at":"2026-02-02T12:06:23Z","stars_delta_30d":32,"open_issues_delta_30d":0},"security_summary":{"status":"no_lockfile","scanner":null,"low_count":0,"high_count":0,"last_scan_at":"2026-07-11T11:43:05.876Z","medium_count":0,"scan_profile":"none","critical_count":0}},"capability_facts":{"scan":{"source":"repo_scan","observed_at":"2026-08-24T18:01:34.364Z"},"languages":{"value":["c++"],"source":"github.language","observed_at":"2026-08-24T18:01:34.364Z"},"license_spdx":{"value":"MIT","source":"github.license","observed_at":"2026-08-24T18:01:34.364Z"}},"decision_facts":{"hosting":null,"pricing":null,"requirements":null,"constraints":null,"when_to_use":["When you have multiple interconnected home devices and want to maximize their combined computing power for LLM inference tasks.","If your use case involves real-time or near-real-time responses where performance acceleration is critical."],"when_not_to_use":["For scenarios with fewer than two available devices, as the framework's capability to distribute and boost performance would be limited.","In professional environments that require strict data privacy controls, due to potential network vulnerabilities among home devices."],"source":"enrich:decision_facts","observed_at":"2026-07-17T08:21:52.218Z"},"constraint_facets":null,"decision_summary":[{"label":"Adopt for","value":"distributed-llama is a C++ framework that leverages multiple home devices for faster large language model inference, under the MIT license."}]}}