{
  "generated_at": "2026-09-16T14:02:18.180Z",
  "total": 16,
  "exported": 16,
  "truncated": false,
  "filters": {
    "hardwareId": null,
    "modelId": "unknown:qwen-qwen2-5-7b-instruct-gguf",
    "quantId": null,
    "frameworkId": null,
    "platform": null,
    "minEvidence": "L0",
    "range": "all",
    "sort": "recent",
    "page": 1,
    "perPage": 50,
    "keyword": null,
    "scope": "all"
  },
  "data_status": "sample",
  "runs": [
    {
      "model": "Qwen2.5-7B-Instruct · 未归一化",
      "quant": "AWQ 4-bit",
      "hardware": "Apple M3 Max · 128G 统一内存",
      "framework": "vLLM 0.6.2",
      "decode_tps": 70.5,
      "prefill_tps": 2466.2,
      "ttft_cold_ms": 2421,
      "vram_peak_gb": 6.1,
      "sample_size": 7,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "待判定",
      "source_platform": "知乎",
      "source_url": "https://zhuanlan.zhihu.com/p/araoai-demo-4",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Qwen2.5-7B-Instruct · 未归一化",
      "quant": "Q4_K_M",
      "hardware": "AMD Radeon RX 7900 XTX · 24G",
      "framework": "mlx-lm 0.18.2",
      "decode_tps": 138.9,
      "prefill_tps": 4699.7,
      "ttft_cold_ms": 2530,
      "vram_peak_gb": 6.4,
      "sample_size": 17,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "稳定",
      "source_platform": "GitHub",
      "source_url": "https://github.com/araoai-demo/benchmarks/issues/1014",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Qwen2.5-7B-Instruct · 未归一化",
      "quant": "Q8_0",
      "hardware": "RX 7800 XT · 未归一化",
      "framework": "mlx-lm 0.18.2",
      "decode_tps": 51,
      "prefill_tps": 1725.1,
      "ttft_cold_ms": 3879,
      "vram_peak_gb": 10.1,
      "sample_size": 3,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "待判定",
      "source_platform": "Reddit",
      "source_url": "https://www.reddit.com/r/araoai_demo/comments/demo15_qwen2-5-7b-instruct-gguf/",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Qwen2.5-7B-Instruct · 未归一化",
      "quant": "AWQ 4-bit",
      "hardware": "AMD Radeon RX 7900 XTX · 24G",
      "framework": "vLLM 0.6.2",
      "decode_tps": 169.3,
      "prefill_tps": 5918.8,
      "ttft_cold_ms": 2421,
      "vram_peak_gb": 6.1,
      "sample_size": 4,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "待判定",
      "source_platform": "HuggingFace",
      "source_url": "https://huggingface.co/araoai-demo/qwen2-5-7b-instruct-gguf-bench/discussions/16",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Qwen2.5-7B-Instruct · 未归一化",
      "quant": "Q5_K_M",
      "hardware": "RTX 3060 · 未归一化",
      "framework": "mlx-lm 0.18.2",
      "decode_tps": 44.6,
      "prefill_tps": 1510.6,
      "ttft_cold_ms": 2822,
      "vram_peak_gb": 7.2,
      "sample_size": 5,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "待判定",
      "source_platform": "V2EX",
      "source_url": "https://www.v2ex.com/t/araoai-demo-17",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Qwen2.5-7B-Instruct · 未归一化",
      "quant": "Q4_K_M",
      "hardware": "M2 Ultra · 未归一化",
      "framework": "ollama 0.5.1",
      "decode_tps": 109.8,
      "prefill_tps": 3672.7,
      "ttft_cold_ms": 2530,
      "vram_peak_gb": 6.4,
      "sample_size": 6,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "待判定",
      "source_platform": "知乎",
      "source_url": "https://zhuanlan.zhihu.com/p/araoai-demo-18",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Qwen2.5-7B-Instruct · 未归一化",
      "quant": "FP16",
      "hardware": "M2 Ultra · 未归一化",
      "framework": "mlx-lm 0.18.2",
      "decode_tps": 34.7,
      "prefill_tps": 1174.9,
      "ttft_cold_ms": 6613,
      "vram_peak_gb": 17.6,
      "sample_size": 11,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "稳定",
      "source_platform": "HuggingFace",
      "source_url": "https://huggingface.co/araoai-demo/qwen2-5-7b-instruct-gguf-bench/discussions/23",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Qwen2.5-7B-Instruct · 未归一化",
      "quant": "AWQ 4-bit",
      "hardware": "RTX 4080 SUPER · 未归一化",
      "framework": "mlx-lm 0.18.2",
      "decode_tps": 113.6,
      "prefill_tps": 3843.3,
      "ttft_cold_ms": 2421,
      "vram_peak_gb": 6.1,
      "sample_size": 14,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "稳定",
      "source_platform": "B站",
      "source_url": "https://www.bilibili.com/video/araoai-demo-26",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Qwen2.5-7B-Instruct · 未归一化",
      "quant": "Q6_K",
      "hardware": "M2 Ultra · 未归一化",
      "framework": "ExLlamaV2 0.2.5",
      "decode_tps": 91,
      "prefill_tps": 3139.1,
      "ttft_cold_ms": 3186,
      "vram_peak_gb": 8.2,
      "sample_size": 15,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "稳定",
      "source_platform": "Arao 自产",
      "source_url": "https://araoai.com/demo/rig-run-27",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Qwen2.5-7B-Instruct · 未归一化",
      "quant": "Q5_K_M",
      "hardware": "Apple M3 Max · 128G 统一内存",
      "framework": "ollama 0.5.1",
      "decode_tps": 47.1,
      "prefill_tps": 1574,
      "ttft_cold_ms": 2822,
      "vram_peak_gb": 7.2,
      "sample_size": 8,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "稳定",
      "source_platform": "GitHub",
      "source_url": "https://github.com/araoai-demo/benchmarks/issues/1035",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Qwen2.5-7B-Instruct · 未归一化",
      "quant": "Q8_0",
      "hardware": "NVIDIA RTX 3090 · 24G",
      "framework": "llama.cpp b6120",
      "decode_tps": 78,
      "prefill_tps": 2652.9,
      "ttft_cold_ms": 3879,
      "vram_peak_gb": 10.1,
      "sample_size": 14,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "稳定",
      "source_platform": "Arao 自产",
      "source_url": "https://araoai.com/demo/rig-run-41",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Qwen2.5-7B-Instruct · 未归一化",
      "quant": "Q4_K_M",
      "hardware": "Arc A770 · 未归一化",
      "framework": "ExLlamaV2 0.2.5",
      "decode_tps": 87.6,
      "prefill_tps": 3021.4,
      "ttft_cold_ms": 2530,
      "vram_peak_gb": 6.4,
      "sample_size": 15,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "稳定",
      "source_platform": "GitHub",
      "source_url": "https://github.com/araoai-demo/benchmarks/issues/1042",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Qwen2.5-7B-Instruct · 未归一化",
      "quant": "AWQ 4-bit",
      "hardware": "NVIDIA RTX 4090 · 24G",
      "framework": "mlx-lm 0.18.2",
      "decode_tps": 155.5,
      "prefill_tps": 5263.7,
      "ttft_cold_ms": 2421,
      "vram_peak_gb": 6.1,
      "sample_size": 17,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "稳定",
      "source_platform": "HuggingFace",
      "source_url": "https://huggingface.co/araoai-demo/qwen2-5-7b-instruct-gguf-bench/discussions/44",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Qwen2.5-7B-Instruct · 未归一化",
      "quant": "Q6_K",
      "hardware": "NVIDIA RTX 3090 · 24G",
      "framework": "ollama 0.5.1",
      "decode_tps": 93.5,
      "prefill_tps": 3125.1,
      "ttft_cold_ms": 3186,
      "vram_peak_gb": 8.2,
      "sample_size": 3,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "待判定",
      "source_platform": "V2EX",
      "source_url": "https://www.v2ex.com/t/araoai-demo-45",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Qwen2.5-7B-Instruct · 未归一化",
      "quant": "AWQ 4-bit",
      "hardware": "RTX 3060 · 未归一化",
      "framework": "mlx-lm 0.18.2",
      "decode_tps": 55.6,
      "prefill_tps": 1879.9,
      "ttft_cold_ms": 2421,
      "vram_peak_gb": 6.1,
      "sample_size": 11,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "稳定",
      "source_platform": "知乎",
      "source_url": "https://zhuanlan.zhihu.com/p/araoai-demo-53",
      "fetched_at": "2026-09-14T00:00:00Z"
    },
    {
      "model": "Qwen2.5-7B-Instruct · 未归一化",
      "quant": "AWQ 4-bit",
      "hardware": "RX 7900 GRE · 未归一化",
      "framework": "ExLlamaV2 0.2.5",
      "decode_tps": 96.1,
      "prefill_tps": 3314.9,
      "ttft_cold_ms": 2421,
      "vram_peak_gb": 6.1,
      "sample_size": 5,
      "evidence": "L0",
      "repro_count": 0,
      "repro_chain": "待复现",
      "status": "待判定",
      "source_platform": "Arao 自产",
      "source_url": "https://araoai.com/demo/rig-run-62",
      "fetched_at": "2026-09-14T00:00:00Z"
    }
  ]
}
