{
  "generated_at": "2026-09-15T11:18:28.604Z",
  "total": 6,
  "exported": 6,
  "truncated": false,
  "filters": {
    "hardwareId": "rtx4090",
    "modelId": null,
    "quantId": null,
    "frameworkId": null,
    "platform": null,
    "minEvidence": "L0",
    "range": "7d",
    "sort": "recent",
    "page": 1,
    "perPage": 50,
    "keyword": null,
    "scope": "all"
  },
  "data_status": "sample",
  "runs": [
    {
      "model": "Qwen3-14B",
      "quant": "Q5_K_M",
      "hardware": "NVIDIA RTX 4090 · 24G",
      "framework": "mlx-lm 0.18.2",
      "decode_tps": 62.5,
      "prefill_tps": 2114.9,
      "ttft_cold_ms": 4863,
      "vram_peak_gb": 12.8,
      "sample_size": 4,
      "evidence": "L1",
      "repro_count": 1,
      "repro_chain": "一致 1 次，偏差 0 次，共 1 格",
      "status": "待判定",
      "source_platform": "Reddit",
      "source_url": "https://www.reddit.com/r/araoai_demo/comments/demo1_qwen3-14b-gguf/",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Mistral-Nemo-Instruct-2407 · 未归一化",
      "quant": "Q5_K_M",
      "hardware": "NVIDIA RTX 4090 · 24G",
      "framework": "llama.cpp b6120",
      "decode_tps": 74.4,
      "prefill_tps": 2529.6,
      "ttft_cold_ms": 4280,
      "vram_peak_gb": 11.2,
      "sample_size": 8,
      "evidence": "L2",
      "repro_count": 3,
      "repro_chain": "一致 3 次，偏差 0 次，共 3 格",
      "status": "稳定",
      "source_platform": "B站",
      "source_url": "https://www.bilibili.com/video/araoai-demo-5",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Mistral-Nemo-Instruct-2407 · 未归一化",
      "quant": "Q5_K_M",
      "hardware": "NVIDIA RTX 4090 · 24G",
      "framework": "mlx-lm 0.18.2",
      "decode_tps": 72.9,
      "prefill_tps": 2467.3,
      "ttft_cold_ms": 4280,
      "vram_peak_gb": 11.2,
      "sample_size": 9,
      "evidence": "L2",
      "repro_count": 3,
      "repro_chain": "一致 3 次，偏差 0 次，共 3 格",
      "status": "稳定",
      "source_platform": "Arao 自产",
      "source_url": "https://araoai.com/demo/rig-run-600",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Mistral-Nemo-Instruct-2407 · 未归一化",
      "quant": "Q4_K_M",
      "hardware": "NVIDIA RTX 4090 · 24G",
      "framework": "ollama 0.5.1",
      "decode_tps": 80.7,
      "prefill_tps": 2699.4,
      "ttft_cold_ms": 3780,
      "vram_peak_gb": 9.8,
      "sample_size": 5,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "待判定",
      "source_platform": "知乎",
      "source_url": "https://zhuanlan.zhihu.com/p/araoai-demo-32",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Qwen2.5-7B-Instruct · 未归一化",
      "quant": "AWQ 4-bit",
      "hardware": "NVIDIA RTX 4090 · 24G",
      "framework": "mlx-lm 0.18.2",
      "decode_tps": 155.5,
      "prefill_tps": 5263.7,
      "ttft_cold_ms": 2421,
      "vram_peak_gb": 6.1,
      "sample_size": 17,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "稳定",
      "source_platform": "HuggingFace",
      "source_url": "https://huggingface.co/araoai-demo/qwen2-5-7b-instruct-gguf-bench/discussions/44",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Mistral-Nemo-Instruct-2407 · 未归一化",
      "quant": "AWQ 4-bit",
      "hardware": "NVIDIA RTX 4090 · 24G",
      "framework": "ollama 0.5.1",
      "decode_tps": 86.1,
      "prefill_tps": 2879.4,
      "ttft_cold_ms": 3592,
      "vram_peak_gb": 9.3,
      "sample_size": 13,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "稳定",
      "source_platform": "GitHub",
      "source_url": "https://github.com/araoai-demo/benchmarks/issues/1070",
      "fetched_at": "2026-09-14T00:00:00Z"
    }
  ]
}
