{ "runId": "run_6b4b2e62-d3fe-4640-8fa9-7c29ed8019a1", "bundleId": "llamacpp-qwen36_35b_q4_k_m.gguf-959317", "status": "verified", "promptTokens": 40960, "completionTokens": 10240, "contextLength": 5120, "harness": { "version": "0.1.22", "gitSha": "50c24e0" }, "runtime": { "name": "llama.cpp", "version": "b8920", "buildFlags": "metal" }, "model": { "displayName": "Qwen3.6-35B-A3B", "format": "gguf", "quant": "q4_k_m", "architecture": "qwen35moe", "source": "DuoNeural/Qwen3.6-35B-A3B-Code-imatrix-GGUF:qwen36_35b_Q4_K_M.gguf", "fileSizeBytes": 21166758144, "lab": { "name": "Qwen", "slug": "qwen" }, "quantizedBy": { "name": "DuoNeural", "slug": "duoneural" } }, "device": { "cpu": "Apple M5 Pro", "cpuCores": 18, "gpu": "Apple M5 Pro", "gpuCores": 20, "gpuCount": 1, "ramGb": 64, "osName": "macOS", "osVersion": "26.5.1" }, "decodeTpsMean": 71.5, "prefillTpsMean": 1532.6, "ttftP50Ms": 2674.12, "idleTpsMean": 11371.8, "peakRssMb": 11416.7, "trialsPassed": 10, "trialsTotal": 10, "runnabilityScore": 0.8092207069614954, "bundleSha256": "cad932daeefcd917f2e3067469bbb4671a73b5cbcb465ff9e172666a8d44c675", "createdAt": "2026-06-15T21:50:26.604Z"}