{ "runId": "run_7d51779a-c695-42be-a932-5db6a394331e", "bundleId": "llamacpp-gemma-4-12b-it-ud-q4_k_xl.gguf-3aa201", "status": "verified", "promptTokens": 40960, "completionTokens": 10240, "contextLength": 5120, "harness": { "version": "0.1.22", "gitSha": "50c24e0" }, "runtime": { "name": "llama.cpp", "version": "b8920", "buildFlags": "metal" }, "model": { "displayName": "Gemma 4 12B IT", "format": "gguf", "quant": "ud-q4_k_xl", "architecture": "gemma4", "source": "unsloth/gemma-4-12b-it-GGUF:gemma-4-12b-it-UD-Q4_K_XL.gguf", "fileSizeBytes": 7366421920, "lab": { "name": "Google", "slug": "google" }, "quantizedBy": { "name": "Unsloth", "slug": "unsloth" } }, "device": { "cpu": "Apple M5 Pro", "cpuCores": 18, "gpu": "Apple M5 Pro", "gpuCores": 20, "gpuCount": 1, "ramGb": 64, "osName": "macOS", "osVersion": "26.5.1" }, "decodeTpsMean": 32.9, "prefillTpsMean": 619.7, "ttftP50Ms": 6609.63, "idleTpsMean": 6047.8, "peakRssMb": 6191.4, "trialsPassed": 10, "trialsTotal": 10, "runnabilityScore": 0.6995314571707589, "bundleSha256": "e714a1f45db8fe904ac5f0911f68bbfb9412b249671ca57336a8dc00e17b2f71", "createdAt": "2026-07-01T15:38:19.336Z"}