{"data":{"slug":"nightmean-ollitert","name":"OlliteRT","tagline":"Android-based local inference server for OpenAI-compatible LLMs","github_url":"https://github.com/NightMean/OlliteRT","owner":"NightMean","repo":"OlliteRT","owner_avatar_url":"https://avatars.githubusercontent.com/u/5726996?v=4","primary_language":"Kotlin","stars":373,"forks":47,"topics":["android","anthropic-api","gemma","home-assistant","kotlin-android","litert","litert-lm","llm","llm-inference","local-llm","on-device-ai","openai-api"],"archived":false,"github_pushed_at":"2026-09-19T23:57:14+00:00","maintenance_label":"Very active","stars_delta_30d":227,"url":"https://www.graphcanon.com/tools/nightmean-ollitert","markdown_url":"https://www.graphcanon.com/tools/nightmean-ollitert.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/nightmean-ollitert","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=nightmean-ollitert","description":"Turn your Android phone into an OpenAI-compatible LLM inference server - Fully local, private and Open Source","homepage_url":null,"license":"Apache-2.0","open_issues":7,"watchers":7,"ai_summary":"OlliteRT allows users to conduct on-device AI inference using their Android phones by transforming them into servers compatible with the OpenAI API protocol.","readme_excerpt":"## Quick Start\n\n1. **Download & install** the APK\n   <br><a href=\"https://github.com/NightMean/ollitert/releases/latest\"><img src=\"assets/Github/Download_Get_it_on_Github.png\" alt=\"Get it on GitHub\" height=\"45\" /></a>\n2. **Download a model** — **Gemma 4 E2B** is recommended for most devices (2.4 GB, runs on 8 GB RAM)\n3. **Start the server** — Tap the Start Server button on the downloaded model card\n4. **Configure your client** — Use the endpoint shown on the Status screen (e.g. `http://PHONE_IP:8000/v1`) with any OpenAI-compatible client — Open WebUI, OpenClaw, Home Assistant, Python, etc. See **[Client Setup](docs/CLIENT_SETUP.md)** for detailed guides.\n\n> [!IMPORTANT]\n> Requires: Android 12+ · **arm64-v8a** device · 6 GB RAM minimum · 8 GB+ recommended for multimodal models (see [model table](#available-models))\n\n---\n\n## License\n\nLicensed under the [Apache License 2.0](LICENSE).","github_created_at":"2026-04-06T18:29:37+00:00","created_at":"2026-07-15T11:02:18.707532+00:00","updated_at":"2026-09-20T05:09:46.057547+00:00","categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"}],"tags":[{"slug":"android","name":"android"},{"slug":"kotlin-android","name":"kotlin-android"},{"slug":"litert-lm","name":"litert-lm"},{"slug":"local-llm","name":"local-llm"},{"slug":"on-device-ai","name":"on-device-ai"},{"slug":"openai-api","name":"openai-api"}],"trust":{"provenance":{"is_fork":false,"github_id":1203116301,"owner_type":"User","methodology":"github_public_v1","parent_repo":null,"near_duplicate_slugs":[]},"computed_at":"2026-09-20T05:09:44.199Z","maintenance":{"label":"Very active","score":96,"methodology":"github_public_v1","releases_90d":0,"days_since_push":0,"last_release_at":"2026-06-06T13:00:39Z","stars_delta_30d":227,"open_issues_delta_30d":-1},"security_summary":{"status":"no_lockfile","scanner":null,"low_count":0,"high_count":0,"last_scan_at":"2026-07-15T11:02:20.313Z","medium_count":0,"scan_profile":"none","critical_count":0}},"capability_facts":{"scan":{"source":"repo_scan","observed_at":"2026-09-20T05:09:45.239Z"},"languages":{"value":["kotlin"],"source":"github.language","observed_at":"2026-09-20T05:09:45.239Z"},"license_spdx":{"value":"Apache-2.0","source":"github.license","observed_at":"2026-09-20T05:09:45.239Z"}},"decision_facts":{"hosting":null,"pricing":null,"requirements":null,"constraints":null,"when_to_use":["Android phones running Android 12+ arm64-v8a are required for deployment","Ideal for users preferring a private on-device setup without cloud dependency"],"when_not_to_use":["Not suitable for devices with less than 6 GB of RAM, more is needed for performance","Avoid if device does not support arm64-v8a architecture or runs Android versions lower than 12"],"source":"enrich:decision_facts","observed_at":"2026-07-17T09:22:59.851Z"},"constraint_facets":null,"decision_summary":[{"label":"Adopt for","value":"OlliteRT lets Android users run local AI inference with OpenAI-like compatibility using their devices"}]}}