{"data":{"slug":"ngxson-wllama","name":"wllama","tagline":"WebAssembly binding for llama.cpp - Enabling on-browser LLM inference","github_url":"https://github.com/ngxson/wllama","owner":"ngxson","repo":"wllama","owner_avatar_url":"https://avatars.githubusercontent.com/u/7702203?v=4","primary_language":"TypeScript","stars":1159,"forks":117,"topics":["llama","llamacpp","llm","wasm","webassembly"],"archived":false,"github_pushed_at":"2026-06-17T17:32:59+00:00","maintenance_label":"Steady","url":"https://www.graphcanon.com/tools/ngxson-wllama","markdown_url":"https://www.graphcanon.com/tools/ngxson-wllama.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/ngxson-wllama","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=ngxson-wllama","description":"WebAssembly binding for llama.cpp - Enabling on-browser LLM inference","homepage_url":"https://huggingface.co/spaces/ngxson/wllama","license":"MIT","open_issues":53,"watchers":12,"ai_summary":"ngxson/wllama is a repository that provides WebAssembly bindings for llama.cpp, allowing for browser-based LLM inference.","readme_excerpt":"# /!\\ IMPORTANT: Require having docker compose installed","github_created_at":"2024-03-13T23:18:27+00:00","created_at":"2026-07-11T10:39:58.868333+00:00","updated_at":"2026-08-07T18:01:51.50255+00:00","categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"}],"tags":[{"slug":"llama","name":"llama"},{"slug":"llamacpp","name":"llamacpp"},{"slug":"llm","name":"llm"},{"slug":"wasm","name":"wasm"},{"slug":"webassembly","name":"webassembly"}],"trust":{"provenance":{"is_fork":false,"github_id":771771112,"owner_type":"User","methodology":"github_public_v1","parent_repo":null,"near_duplicate_slugs":[]},"computed_at":"2026-08-07T18:01:49.640Z","maintenance":{"label":"Steady","score":60,"methodology":"github_public_v1","releases_90d":8,"days_since_push":51,"last_release_at":"2026-06-15T22:38:12Z"},"security_summary":{"status":"findings","scanner":"osv@v1","low_count":9,"high_count":0,"last_scan_at":"2026-07-11T10:40:01.063Z","medium_count":0,"scan_profile":"deps","critical_count":0}},"capability_facts":{"mcp":{"source":"repo_scan","observed_at":"2026-08-07T18:01:50.098Z","server_manifest":false},"scan":{"source":"repo_scan","observed_at":"2026-08-07T18:01:50.098Z"},"has_cli":{"value":true,"source":"package.json:bin|scripts","observed_at":"2026-08-07T18:01:50.098Z"},"languages":{"value":["typescript","javascript"],"source":"github.language+package.json","observed_at":"2026-08-07T18:01:50.098Z"},"license_spdx":{"value":"MIT","source":"github.license","observed_at":"2026-08-07T18:01:50.098Z"}},"decision_facts":{"hosting":null,"pricing":null,"requirements":null,"constraints":null,"when_to_use":["Need browser-based LLM inference directly through WebAssembly.","Prefer TypeScript environment for developing AI web applications."],"when_not_to_use":["Require direct native execution speed benefits unavailable in a WebAssembly context.","Developing server-side applications without the need for client-side inference capabilities."],"source":"enrich:decision_facts","observed_at":"2026-07-12T11:16:06.271Z"},"constraint_facets":null,"decision_summary":[{"label":"Adopt for","value":"WebAssembly bindings for browser-based inference of llama.cpp."}]}}