{
  "generated_at": "2026-09-15T11:55:11.656Z",
  "total": 2,
  "exported": 2,
  "truncated": false,
  "filters": {
    "hardwareId": "unknown:rtx-4080-super",
    "modelId": null,
    "quantId": null,
    "frameworkId": "llamacpp",
    "platform": null,
    "minEvidence": "L0",
    "range": "all",
    "sort": "recent",
    "page": 1,
    "perPage": 50,
    "keyword": null,
    "scope": "all"
  },
  "data_status": "sample",
  "runs": [
    {
      "model": "Qwen3-14B",
      "quant": "Q6_K",
      "hardware": "RTX 4080 SUPER · 未归一化",
      "framework": "llama.cpp b6120",
      "decode_tps": 39.5,
      "prefill_tps": 1343.3,
      "ttft_cold_ms": 5592,
      "vram_peak_gb": 14.8,
      "sample_size": 9,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "稳定",
      "source_platform": "Reddit",
      "source_url": "https://www.reddit.com/r/araoai_demo/comments/demo36_qwen3-14b-gguf/",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "phi-4 · 未归一化",
      "quant": "AWQ 4-bit",
      "hardware": "RTX 4080 SUPER · 未归一化",
      "framework": "llama.cpp b6120",
      "decode_tps": 57.9,
      "prefill_tps": 1970.1,
      "ttft_cold_ms": 4061,
      "vram_peak_gb": 10.6,
      "sample_size": 14,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "稳定",
      "source_platform": "GitHub",
      "source_url": "https://github.com/araoai-demo/benchmarks/issues/1056",
      "fetched_at": "2026-09-14T00:00:00Z"
    }
  ]
}
