~/llamay $ # what the install line actually does, pointed at a scratch dir ~/llamay $ LLAMAY_INSTALL_DIR=build/demo/work/bin LLAMAY_NO_SERVICE=1 \ > sh packaging/install.sh llamay 0.1.58 — darwin/arm64 (gpu) downloading llamay_0.1.58_darwin_arm64_metal.tar.gz checksum ok installed build/demo/work/bin/llamay run: llamay serve ~/llamay $ # it verified a signature-grade checksum before it wrote anything ~/llamay $ build/demo/work/bin/llamay version llamay 0.1.58 build go1.26.8 darwin/arm64, cgo on, tags metal,accelerate revision 8e08433a6e86 kernels neon+i8mm gpu metal built in; metal:Apple M4 formats 23: F32, F16, Q4_0, Q4_1, Q5_0, Q5_1, Q8_0, Q8_1, Q2_K, Q3_K, Q4_K, Q5_K, Q6_ K, Q8_K, IQ2_XXS, IQ2_XS, IQ3_XXS, IQ4_NL, IQ3_S, IQ2_S, IQ4_XS, BF16, MXFP4 architectures 17 entries, 39 names: azmx (azmx-one, llama, mistral), qwen2 (qwen3), gemma, gemma2, gemma3 (gemma3_text), azmx-code (azmx_code, azmxcode), gpt-oss (gptoss, gpt_oss, ope nai-moe), phi3 (phi3.5), stablelm (stablelm2), gpt2, falcon (rw), gptneox (gpt_neox, gpt-neo x, neox), phi2, qwen3moe (qwen3_moe), qwen2moe (qwen2_moe), azmx-ocr (ocr, ctc), azmx-embed (bert, embedding)