{ "runId": "run_c8ab71bc-977d-4286-bc7e-7bf589404db6", "bundleId": "llamacpp-qwen3.8-27b-ud-q4_k_s.gguf-ac041f", "status": "verified", "promptTokens": 40960, "completionTokens": 10240, "contextLength": 5120, "harness": { "version": "0.1.21", "gitSha": "159b74142" }, "runtime": { "name": "llama.cpp", "version": "0.3.0-dev", "buildFlags": null }, "model": { "displayName": "Qwen3.8-27B", "format": "gguf", "quant": "q4_k_s", "architecture": "qwen35", "source": "unsloth/Qwen3.8-27B-GGUF:Qwen3.8-27B-UD-Q4_K_S.gguf", "fileSizeBytes": 15358213024, "lab": { "name": "Qwen", "slug": "qwen" }, "quantizedBy": { "name": "Unsloth", "slug": "unsloth" } }, "device": { "cpu": "Apple M3 Max", "cpuCores": 16, "gpu": "Apple M3 Max", "gpuCores": 40, "gpuCount": 1, "ramGb": 128, "osName": "macOS", "osVersion": "26.6.2" }, "decodeTpsMean": 12.6, "prefillTpsMean": 170.9, "ttftP50Ms": 22842.36, "idleTpsMean": 12247, "peakRssMb": 17034.4, "trialsPassed": 10, "trialsTotal": 10, "runnabilityScore": 0.4904269321986607, "bundleSha256": "6ef1765cf379bba038eb4dbaf9096b8019c4395381014d879da0c0048578c769", "createdAt": "2026-09-02T22:32:31.973Z"}