{"data":{"slug":"jundot-omlx","name":"omlx","tagline":"LLM inference server with continuous batching and SSD caching for Apple Silicon","github_url":"https://github.com/jundot/omlx","owner":"jundot","repo":"omlx","owner_avatar_url":"https://avatars.githubusercontent.com/u/64250138?v=4","primary_language":"Python","stars":21934,"forks":1899,"topics":["apple-silicon","inference-server","llm","macos","mlx","openai-api"],"archived":false,"github_pushed_at":"2026-09-20T02:23:26+00:00","maintenance_label":"Very active","stars_delta_30d":3255,"url":"https://www.graphcanon.com/tools/jundot-omlx","markdown_url":"https://www.graphcanon.com/tools/jundot-omlx.md","api_url":"https://www.graphcanon.com/api/graphcanon/tools/jundot-omlx","graph_url":"https://www.graphcanon.com/api/graphcanon/graph?tool=jundot-omlx","description":"LLM inference server with continuous batching & SSD caching for Apple Silicon — managed from the macOS menu bar","homepage_url":"https://omlx.ai","license":"Apache-2.0","open_issues":1407,"watchers":115,"ai_summary":"omlx is an LLM inference server designed to run on Apple Silicon that supports continuous batching and uses SSD caching. The tool can be managed via the macOS menu bar or through Homebrew installation.","readme_excerpt":"# Managed background server (macOS app or Homebrew install)\nomlx start\nomlx stop\nomlx restart","github_created_at":"2026-02-13T14:13:27+00:00","created_at":"2026-07-15T11:19:09.39734+00:00","updated_at":"2026-09-20T05:19:11.240208+00:00","categories":[{"slug":"inference-serving","name":"Inference & Serving","url":"https://www.graphcanon.com/categories/inference-serving","markdown_url":"https://www.graphcanon.com/categories/inference-serving.md","api_url":"https://www.graphcanon.com/api/graphcanon/categories/inference-serving"}],"tags":[{"slug":"apple-silicon","name":"apple-silicon"},{"slug":"inference-server","name":"inference-server"},{"slug":"llm","name":"llm"},{"slug":"macos","name":"macos"}],"trust":{"provenance":{"is_fork":false,"github_id":1157171418,"owner_type":"User","methodology":"github_public_v1","parent_repo":null,"near_duplicate_slugs":[]},"computed_at":"2026-09-20T05:19:08.694Z","maintenance":{"label":"Very active","score":96,"methodology":"github_public_v1","releases_90d":30,"days_since_push":0,"last_release_at":"2026-09-18T06:42:55Z","stars_delta_30d":3255,"open_issues_delta_30d":552},"security_summary":{"status":"no_lockfile","scanner":null,"low_count":0,"high_count":0,"last_scan_at":"2026-07-15T11:19:10.652Z","medium_count":0,"scan_profile":"none","critical_count":0}},"capability_facts":{"scan":{"source":"repo_scan","observed_at":"2026-09-20T05:19:09.695Z"},"has_cli":{"value":true,"source":"pyproject.toml:[project.scripts]","observed_at":"2026-09-20T05:19:09.695Z"},"languages":{"value":["python"],"source":"github.language+pyproject.toml","observed_at":"2026-09-20T05:19:09.695Z"},"license_spdx":{"value":"Apache-2.0","source":"github.license","observed_at":"2026-09-20T05:19:09.695Z"}},"decision_facts":{"hosting":null,"pricing":null,"requirements":{"min_ram_gb":null,"requires_docker":false},"constraints":{"min_ram_gb":null,"requires_docker":false},"when_to_use":["If your primary computing environment is based on Apple Silicon devices, omlx offers optimized performance for running large language model inferences.","When you require an inference server that supports continuous batching and leverages SSD caching for quicker service responses.","You need a seamless control interface via the macOS menu bar or prefer the ease of Homebrew installation for server management."],"when_not_to_use":["If your development infrastructure relies on non-Apple Silicon hardware, omlx's specific optimizations will not be as beneficial.","Teams that require cross-platform compatibility or run servers predominantly on non-macOS operating systems should consider alternatives with broader support.","For environments where direct control through the menu bar is not practical or desired."],"source":"enrich:decision_facts","observed_at":"2026-07-17T12:37:29.410Z"},"constraint_facets":{"min_ram_gb":null,"requires_docker":false},"decision_summary":[{"label":"Adopt for","value":"omlx is an LLM inference server tailored for Apple Silicon that emphasizes continuous batching and SSD caching capabilities, accessible through macOS menu bar control or Homebrew installation."},{"label":"License detail","value":"omlx is available under the Apache License, Version 2.0 (Apache-2.0), a permissive free software license."}]}}