{
  "model_tag": "gemma4:26b",
  "arm": "one-dynboost",
  "box_class": "gpu-5090-laptop-24g",
  "mode": "concurrency",
  "target_uuid": "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff",
  "cards": [
    "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff"
  ],
  "card_labels": {
    "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": "soldered (mobile; no socket \u2014 link width is traced under load, never assumed) at 00000000:01:00.0 - vendor unread - GPU-edff232c"
  },
  "card_identity": {
    "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": {
      "index": null,
      "bus_id": "00000000:01:00.0",
      "pci_sub_device_id": null,
      "vendor": null,
      "board_per_operator": null,
      "vbios": null,
      "serial": null,
      "pcie_link_width_max": null,
      "pcie_link_width_negotiated": null,
      "memory_total_mib": null,
      "memory_bus_width_bits": null,
      "clocks_max_sm_mhz": null,
      "clocks_max_memory_mhz": null,
      "power_default_limit_w": null,
      "power_min_limit_w": null,
      "power_max_limit_w": null,
      "power_limit_at_bench_start_w": null,
      "note": "the single NVIDIA GeForce RTX 5090 Laptop GPU 24 GB SOLDERED to this machine's board at 00000000:01:00.0. There is no seat, no partner and no card to swap. \u26a0 THE TRAP ON THIS BOX IS NOT A STALE BOOT UNIT, IT IS A RUNNING DAEMON: nvidia-powerd.service (NVIDIA Dynamic Boost) floats this board's power limit between its default and its maximum against the CPU's draw, continuously and without asking, and it has been running since 2026-09-08. The lesson every leg of this ladder carries \u2014 a cap read once at the start of a night is not a cap \u2014 is true here in its strongest form: such a reading is a sample of a moving signal. Every arm reads the limit back in the same call as its own data, every stage reads it again at its close, and a stage that finds it moved STOPS rather than relabelling its file. The second trap is a CLOCK LOCK: ai-perf.service runs `nvidia-smi -lgc 1200,2550` at boot on this box and the cards that produced the rows this bench compares against ran unlocked, so a floored clock changes what a power cap means and every record carries clocks.sm min and max so the confound is visible rather than implied."
    }
  },
  "num_ctx": 4096,
  "instance_env": {
    "kind": "instance-env-receipt",
    "box_class": "gpu-5090-laptop-24g",
    "target_uuid": "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff",
    "written_utc": "2026-09-21T20:49:14Z",
    "unit": "bench-5090laptop-one",
    "shape": "one",
    "port": 11470,
    "base_url": "http://127.0.0.1:11470",
    "cuda_visible_devices": "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff",
    "OLLAMA_FLASH_ATTENTION": "1",
    "OLLAMA_KV_CACHE_TYPE": "q8_0",
    "OLLAMA_CONTEXT_LENGTH": 32768,
    "OLLAMA_<parallel-requests>": 4,
    "OLLAMA_MAX_LOADED_MODELS": 1,
    "OLLAMA_KEEP_ALIVE": "0 (the server default; every bench request sends its own keep_alive, which wins, and every arm unloads on the way out)",
    "models_dir": "/usr/share/ollama/.ollama/models",
    "models_dir_writable_by_this_user": false,
    "api_version": {
      "version": "0.32.13"
    },
    "ollama_binary": "/workshop/bench-laptop-5090-2026-09-21/ollama-0.32.13/bin/ollama",
    "ollama_client_version": "0.32.13",
    "ollama_version_pin": "0.32.13",
    "disclosed_tenants": "nomic-embed-text:latest on the SYSTEM ollama (:<the runtime's default port>, OLLAMA_KEEP_ALIVE=-1, ~323 MB VRAM) \u2014 disclosed, never unloaded",
    "power_and_clocks_at_start_csv": "[N/A], 150.00 W, 95.00 W, 175.00 W, 1822 MHz, 14001 MHz, 3090 MHz",
    "power_and_clocks_at_start_fields": "power.limit,enforced.power.limit,power.default_limit,power.max_limit,clocks.sm,clocks.mem,clocks.max.sm",
    "nvidia_powerd_at_start": "active",
    "clock_lock_unit_at_start": "active",
    "clock_lock_unit_name": "ai-perf.service",
    "clock_lock_range_declared": "1200,2550",
    "clock_lock_unit_state_is_not_evidence": "ai-perf.service is Type=oneshot and reads 'active (exited)' for the whole uptime whatever happens to the card afterwards. The lock's real state is read from the CARD (clocks.sm against the lock's floor) and is recorded in this leg's clock-lock receipt.",
    "power_limits_at_start_w": "0, [N/A], 150.00 W;",
    "persistence_mode_at_start": "0, Enabled;",
    "pcie_link_width_at_start": "0, 8;"
  },
  "num_gpu_option": null,
  "num_gpu_policy": "layers left to ollama's own planner (`auto`)",
  "num_gpu_note": "This arm MUST use the same layer policy as the arm it is compared against. On the pair of 3080s, left to the planner, this model spilled a quarter of itself to host RAM and the concurrency figures would have measured that spill rather than the cards; on one 24 GB card the planner is expected to place every layer, and the placement actually measured is recorded per run either way.",
  "memory_temp_support": {
    "checked_utc": "2026-09-21T20:49:14Z",
    "paths": {
      "nvidia-smi --query-gpu=temperature.memory": {
        "raw": "0, GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff, N/A",
        "readable": false
      },
      "nvidia-smi -q -d TEMPERATURE": {
        "lines": [
          "GPU Target Temperature                         : 87 C",
          "Memory Current Temp                            : N/A",
          "Memory Max Operating T.Limit Temp              : N/A"
        ],
        "readable": false
      },
      "NVML NVML_FI_DEV_MEMORY_TEMP": {
        "available": true,
        "field_id": 82,
        "cards": {
          "0": {
            "call_rc": 0,
            "field_rc": 3,
            "supported": false,
            "not_supported": true,
            "value_c": null
          }
        }
      }
    },
    "readable": false,
    "verdict": "the memory die's temperature is NOT readable on these cards through any of the three paths asked, so the memory-temperature stop condition could not arm and NO memory temperature is reported anywhere in this bench. The core temperature stop and the driver's own thermal-slowdown reasons are the thermal instrument instead."
  },
  "ups_before_arm": {
    "ups": "pr1500@localhost",
    "read_utc": "2026-09-21T20:49:14Z",
    "available": false,
    "error": "no upsc on this box"
  },
  "parallel_kv_note": "ollama allocates a KV cache PER PARALLEL SLOT, so a level of 4 streams at num_ctx N costs about four times the KV of one stream at N. The window this arm ran at is recorded above, and it is not automatically the largest window the card holds for a single stream.",
  "base_url": "http://127.0.0.1:11470",
  "started_utc": "2026-09-21T20:49:14Z",
  "prompt_sha256": "90eedd0c53f9554ae3837674504fcb7090013d9a432a653352183c0f25a7ce5c",
  "num_predict": 256,
  "trials_per_level": 3,
  "levels": [
    {
      "level": 1,
      "contention_gate": {
        "attempts": [
          {
            "window_s": 10.0,
            "samples": 20,
            "cards": {
              "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": {
                "utilization_pct": {
                  "min": 0.0,
                  "median": 0.0,
                  "max": 92.0,
                  "mean": 32.2,
                  "n": 20
                },
                "power_w": {
                  "min": 22.67,
                  "median": 22.73,
                  "max": 146.03,
                  "mean": 32.44,
                  "n": 20
                },
                "card_seen": true
              }
            },
            "max_mean_util_pct": 32.2,
            "all_cards_seen": true,
            "attempt": 1,
            "verdict": "busy: busiest card mean 32.2% > 5.0% -- waiting"
          },
          {
            "window_s": 10.0,
            "samples": 20,
            "cards": {
              "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": {
                "utilization_pct": {
                  "min": 0.0,
                  "median": 0.0,
                  "max": 0.0,
                  "mean": 0.0,
                  "n": 20
                },
                "power_w": {
                  "min": 7.03,
                  "median": 22.62,
                  "max": 22.66,
                  "mean": 20.72,
                  "n": 20
                },
                "card_seen": true
              }
            },
            "max_mean_util_pct": 0.0,
            "all_cards_seen": true,
            "attempt": 2,
            "verdict": "quiet: busiest card mean 0.0% <= 5.0%"
          }
        ],
        "passed": true,
        "cards": [
          "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff"
        ]
      },
      "trials": [
        {
          "level": 1,
          "wall_s": 2.0361,
          "streams": [
            {
              "ttft_ms": 386.95,
              "first_token_was_thinking": false,
              "wall_s": 2.0355,
              "eval_count": 256,
              "eval_duration_ns": 1645799000,
              "decode_tok_s": 155.548,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 79396000,
              "prefill_tok_s": 6750.97,
              "load_duration_ns": 301158101,
              "total_duration_ns": 2032103320,
              "streamed_chunks": 256,
              "response_chars": 1044,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 154.685,
              "decode_tok_s_wall_gross": 125.77,
              "num_ctx_option": 4096,
              "stream": 1
            }
          ],
          "streams_ok": 1,
          "streams_failed": 0,
          "stream_errors": [],
          "trial_void": false,
          "per_stream_decode_tok_s": [
            155.548
          ],
          "per_stream_ttft_ms": [
            386.95
          ],
          "aggregate_sum_tok_s": 155.548,
          "aggregate_wall_tok_s": 125.73,
          "tokens_decoded_total": 256,
          "per_stream_decode": {
            "min": 155.548,
            "median": 155.548,
            "max": 155.548,
            "mean": 155.548,
            "n": 1
          },
          "per_stream_ttft": {
            "min": 386.95,
            "median": 386.95,
            "max": 386.95,
            "mean": 386.95,
            "n": 1
          },
          "power": {
            "idle_baseline": {
              "seconds": 3.0,
              "samples": 6,
              "power_w": {
                "min": 6.55,
                "median": 6.59,
                "max": 7.03,
                "mean": 6.72,
                "n": 6
              },
              "utilization_pct": {
                "min": 0.0,
                "median": 0.0,
                "max": 0.0,
                "mean": 0.0,
                "n": 6
              },
              "memory_used_mib": {
                "min": 19176.0,
                "median": 19176.0,
                "max": 19176.0,
                "mean": 19176.0,
                "n": 6
              },
              "power_w_per_card": {
                "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": 6.72
              },
              "memory_used_mib_per_card": {
                "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": 19176.0
              },
              "power_w_all_cards": 6.72,
              "ups_whole_box": {
                "load_pct": null,
                "realpower_w": null,
                "realpower_nominal_w": null,
                "status": null,
                "note": "the UPS's reading of EVERYTHING it feeds -- the whole box, not the cards. The card figures beside it are board power from nvidia-smi. Neither is the other."
              }
            },
            "mean_w": 148.9,
            "max_w": 151.39,
            "over_idle_w": 142.18,
            "samples": 8,
            "per_gpu": {
              "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": {
                "index": "0",
                "name": "NVIDIA GeForce RTX 5090 Laptop GPU",
                "power_w": {
                  "min": 6.45,
                  "median": 147.89,
                  "max": 151.39,
                  "mean": 113.47,
                  "n": 8
                },
                "power_limit_w": 140.0,
                "power_limit_field": "enforced.power.limit",
                "power_limit_enforced_w": 140.0,
                "cap": "dynamic-boost",
                "power_posture": "dynamic-boost",
                "power_limit_enforced_stats_w": {
                  "min": 140.0,
                  "median": 150.0,
                  "max": 150.0,
                  "mean": 147.5,
                  "n": 8
                },
                "power_posture_note": "dynamic-boost -- nvidia-powerd running and no cap set: the board floats between its power.default_limit and its power.max_limit against the CPU's draw, DURING the scored run. `power_limit_enforced_stats_w` is that limit reduced over this run's own 2 Hz trace and is the only honest cap cell; `power_limit_enforced_w` is the FIRST sample and is kept only so the two can be compared. A single wattage in a cap column for this posture is a fabrication.",
                "cap_cell": "140-150 W (median 150, n=8)",
                "memory_used_mib": {
                  "min": 19176.0,
                  "median": 19176.0,
                  "max": 19176.0,
                  "mean": 19176.0,
                  "n": 8
                },
                "utilization_pct": {
                  "min": 0.0,
                  "median": 92.0,
                  "max": 93.0,
                  "mean": 57.6,
                  "n": 8
                },
                "temperature_c": {
                  "min": 35.0,
                  "median": 42.5,
                  "max": 44.0,
                  "mean": 40.9,
                  "n": 8
                },
                "memory_temperature_c": null,
                "fan_pct": null,
                "pcie_width_current": {
                  "min": 8.0,
                  "median": 8.0,
                  "max": 8.0,
                  "mean": 8.0,
                  "n": 8
                },
                "clocks_sm_mhz": {
                  "min": 180.0,
                  "median": 2246.0,
                  "max": 2370.0,
                  "mean": 1532.6,
                  "n": 8
                },
                "clocks_mem_mhz": {
                  "min": 405.0,
                  "median": 14001.0,
                  "max": 14001.0,
                  "mean": 10602.0,
                  "n": 8
                },
                "throttle_sw_power_cap_fraction": 0.625,
                "throttle_hw_slowdown_fraction": 0.0,
                "throttle_sw_thermal_fraction": 0.0,
                "throttle_hw_thermal_fraction": 0.0,
                "n": 8,
                "busy": {
                  "index": "0",
                  "name": "NVIDIA GeForce RTX 5090 Laptop GPU",
                  "power_w": {
                    "min": 146.3,
                    "median": 148.63,
                    "max": 151.39,
                    "mean": 148.9,
                    "n": 5
                  },
                  "power_limit_w": 150.0,
                  "power_limit_field": "enforced.power.limit",
                  "power_limit_enforced_w": 150.0,
                  "cap": "dynamic-boost",
                  "power_posture": "dynamic-boost",
                  "power_limit_enforced_stats_w": {
                    "min": 150.0,
                    "median": 150.0,
                    "max": 150.0,
                    "mean": 150.0,
                    "n": 5
                  },
                  "power_posture_note": "dynamic-boost -- nvidia-powerd running and no cap set: the board floats between its power.default_limit and its power.max_limit against the CPU's draw, DURING the scored run. `power_limit_enforced_stats_w` is that limit reduced over this run's own 2 Hz trace and is the only honest cap cell; `power_limit_enforced_w` is the FIRST sample and is kept only so the two can be compared. A single wattage in a cap column for this posture is a fabrication.",
                  "cap_cell": "150 W, flat",
                  "cap_cell_note": "the enforced limit read 150 W on every one of 5 samples of this run's 2 Hz trace. Under dynamic-boost that is a FINDING, not a range.",
                  "memory_used_mib": {
                    "min": 19176.0,
                    "median": 19176.0,
                    "max": 19176.0,
                    "mean": 19176.0,
                    "n": 5
                  },
                  "utilization_pct": {
                    "min": 92.0,
                    "median": 92.0,
                    "max": 93.0,
                    "mean": 92.2,
                    "n": 5
                  },
                  "temperature_c": {
                    "min": 42.0,
                    "median": 43.0,
                    "max": 44.0,
                    "mean": 43.2,
                    "n": 5
                  },
                  "memory_temperature_c": null,
                  "fan_pct": null,
                  "pcie_width_current": {
                    "min": 8.0,
                    "median": 8.0,
                    "max": 8.0,
                    "mean": 8.0,
                    "n": 5
                  },
                  "clocks_sm_mhz": {
                    "min": 2242.0,
                    "median": 2340.0,
                    "max": 2370.0,
                    "mean": 2312.8,
                    "n": 5
                  },
                  "clocks_mem_mhz": {
                    "min": 14001.0,
                    "median": 14001.0,
                    "max": 14001.0,
                    "mean": 14001.0,
                    "n": 5
                  },
                  "throttle_sw_power_cap_fraction": 1.0,
                  "throttle_hw_slowdown_fraction": 0.0,
                  "throttle_sw_thermal_fraction": 0.0,
                  "throttle_hw_thermal_fraction": 0.0,
                  "n": 5
                },
                "busy_samples": 5
              }
            }
          },
          "energy_j_per_1k_tokens": 1184.28,
          "thermal": {
            "rule_version": 2,
            "measured_over": "samples where the GPU was busy (utilisation > 0)",
            "busy_samples": 5,
            "temp_max_c": 44.0,
            "temp_min_c": 42.0,
            "temp_rise_c": 2.0,
            "memory_temp_max_c": null,
            "memory_temp_median_c": null,
            "memory_temp_readable": false,
            "pcie_width_min": 8.0,
            "pcie_width_max": 8.0,
            "pcie_width_median": 8.0,
            "fan_mean_pct": null,
            "fan_max_pct": null,
            "clock_floor_mhz": 2242.0,
            "clock_max_mhz": 2370.0,
            "clock_median_mhz": 2340.0,
            "mem_clock_floor_mhz": 14001.0,
            "mem_clock_median_mhz": 14001.0,
            "mem_clock_max_mhz": 14001.0,
            "mem_clock_held": true,
            "memory_bus_width_bits": 256,
            "memory_technology": "GDDR7",
            "memory_bits_per_clock": null,
            "peak_mem_bandwidth_gbs_at_median_clock": null,
            "peak_mem_bandwidth_gbs_at_floor_clock": null,
            "peak_mem_bandwidth_gbs_at_max_clock": null,
            "peak_mem_bandwidth_withheld_because": "this harness's peak-bandwidth formula assumes two bits per memory clock per pin, which is exact for GDDR6 and GDDR6X and is checked against three published bandwidths. It has no receipt for the bits-per-clock of this board's memory, so the DERIVED bytes-per-second ceiling is withheld rather than computed from an unverified constant. The memory clock itself is a reading and is published unchanged.",
            "peak_mem_bandwidth_formula": "GB/s = clocks.mem MHz x 1e6 x <bits-per-clock for THIS board's memory> x (bus bits / 8) / 1e9; a CEILING at the clock that was observed, never an achieved rate. The bits-per-clock is registered per board and is NOT defaulted.",
            "clock_dropped": true,
            "power_limit_w": 140.0,
            "power_mean_w": 148.9,
            "power_pct_of_limit": 106.4,
            "power_pinned_at_limit": true,
            "sw_power_cap_active_fraction": 1.0,
            "hw_slowdown_active_fraction": 0.0,
            "sw_thermal_active_fraction": 0.0,
            "hw_thermal_active_fraction": 0.0,
            "at_throttling_temperature": false,
            "stop_core_temp_c": 83.0,
            "stop_memory_temp_c": 100.0,
            "hit_core_temp_stop": false,
            "hit_memory_temp_stop": false,
            "throttle_temperature_c": 80.0,
            "temp_rising_within_run": false,
            "verdict": "the POWER CAP doing its job: the SM clock fell while draw sat at 106.4% of the 140 W cap, at 44 C -- well below any throttling temperature"
          }
        },
        {
          "level": 1,
          "wall_s": 1.9794,
          "streams": [
            {
              "ttft_ms": 329.71,
              "first_token_was_thinking": false,
              "wall_s": 1.9787,
              "eval_count": 256,
              "eval_duration_ns": 1645189000,
              "decode_tok_s": 155.605,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 72244000,
              "prefill_tok_s": 7419.301,
              "load_duration_ns": 252061617,
              "total_duration_ns": 1974351616,
              "streamed_chunks": 256,
              "response_chars": 1044,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 154.636,
              "decode_tok_s_wall_gross": 129.376,
              "num_ctx_option": 4096,
              "stream": 1
            }
          ],
          "streams_ok": 1,
          "streams_failed": 0,
          "stream_errors": [],
          "trial_void": false,
          "per_stream_decode_tok_s": [
            155.605
          ],
          "per_stream_ttft_ms": [
            329.71
          ],
          "aggregate_sum_tok_s": 155.605,
          "aggregate_wall_tok_s": 129.333,
          "tokens_decoded_total": 256,
          "per_stream_decode": {
            "min": 155.605,
            "median": 155.605,
            "max": 155.605,
            "mean": 155.605,
            "n": 1
          },
          "per_stream_ttft": {
            "min": 329.71,
            "median": 329.71,
            "max": 329.71,
            "mean": 329.71,
            "n": 1
          },
          "power": {
            "idle_baseline": {
              "seconds": 3.0,
              "samples": 6,
              "power_w": {
                "min": 22.71,
                "median": 30.32,
                "max": 146.24,
                "mean": 55.55,
                "n": 6
              },
              "utilization_pct": {
                "min": 0.0,
                "median": 0.0,
                "max": 93.0,
                "mean": 15.5,
                "n": 6
              },
              "memory_used_mib": {
                "min": 19176.0,
                "median": 19176.0,
                "max": 19176.0,
                "mean": 19176.0,
                "n": 6
              },
              "power_w_per_card": {
                "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": 55.55
              },
              "memory_used_mib_per_card": {
                "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": 19176.0
              },
              "power_w_all_cards": 55.55,
              "ups_whole_box": {
                "load_pct": null,
                "realpower_w": null,
                "realpower_nominal_w": null,
                "status": null,
                "note": "the UPS's reading of EVERYTHING it feeds -- the whole box, not the cards. The card figures beside it are board power from nvidia-smi. Neither is the other."
              }
            },
            "mean_w": 146.96,
            "max_w": 148.33,
            "over_idle_w": 91.41,
            "samples": 8,
            "per_gpu": {
              "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": {
                "index": "0",
                "name": "NVIDIA GeForce RTX 5090 Laptop GPU",
                "power_w": {
                  "min": 22.62,
                  "median": 82.28,
                  "max": 148.33,
                  "mean": 86.36,
                  "n": 8
                },
                "power_limit_w": 150.0,
                "power_limit_field": "enforced.power.limit",
                "power_limit_enforced_w": 150.0,
                "cap": "dynamic-boost",
                "power_posture": "dynamic-boost",
                "power_limit_enforced_stats_w": {
                  "min": 150.0,
                  "median": 150.0,
                  "max": 150.0,
                  "mean": 150.0,
                  "n": 8
                },
                "power_posture_note": "dynamic-boost -- nvidia-powerd running and no cap set: the board floats between its power.default_limit and its power.max_limit against the CPU's draw, DURING the scored run. `power_limit_enforced_stats_w` is that limit reduced over this run's own 2 Hz trace and is the only honest cap cell; `power_limit_enforced_w` is the FIRST sample and is kept only so the two can be compared. A single wattage in a cap column for this posture is a fabrication.",
                "cap_cell": "150 W, flat",
                "cap_cell_note": "the enforced limit read 150 W on every one of 8 samples of this run's 2 Hz trace. Under dynamic-boost that is a FINDING, not a range.",
                "memory_used_mib": {
                  "min": 19176.0,
                  "median": 19176.0,
                  "max": 19176.0,
                  "mean": 19176.0,
                  "n": 8
                },
                "utilization_pct": {
                  "min": 0.0,
                  "median": 0.0,
                  "max": 93.0,
                  "mean": 34.9,
                  "n": 8
                },
                "temperature_c": {
                  "min": 36.0,
                  "median": 43.5,
                  "max": 45.0,
                  "mean": 41.8,
                  "n": 8
                },
                "memory_temperature_c": null,
                "fan_pct": null,
                "pcie_width_current": {
                  "min": 8.0,
                  "median": 8.0,
                  "max": 8.0,
                  "mean": 8.0,
                  "n": 8
                },
                "clocks_sm_mhz": {
                  "min": 1590.0,
                  "median": 2280.0,
                  "max": 2362.0,
                  "mean": 2128.9,
                  "n": 8
                },
                "clocks_mem_mhz": {
                  "min": 14001.0,
                  "median": 14001.0,
                  "max": 14001.0,
                  "mean": 14001.0,
                  "n": 8
                },
                "throttle_sw_power_cap_fraction": 0.625,
                "throttle_hw_slowdown_fraction": 0.0,
                "throttle_sw_thermal_fraction": 0.0,
                "throttle_hw_thermal_fraction": 0.0,
                "n": 8,
                "busy": {
                  "index": "0",
                  "name": "NVIDIA GeForce RTX 5090 Laptop GPU",
                  "power_w": {
                    "min": 146.04,
                    "median": 146.51,
                    "max": 148.33,
                    "mean": 146.96,
                    "n": 3
                  },
                  "power_limit_w": 150.0,
                  "power_limit_field": "enforced.power.limit",
                  "power_limit_enforced_w": 150.0,
                  "cap": "dynamic-boost",
                  "power_posture": "dynamic-boost",
                  "power_limit_enforced_stats_w": {
                    "min": 150.0,
                    "median": 150.0,
                    "max": 150.0,
                    "mean": 150.0,
                    "n": 3
                  },
                  "power_posture_note": "dynamic-boost -- nvidia-powerd running and no cap set: the board floats between its power.default_limit and its power.max_limit against the CPU's draw, DURING the scored run. `power_limit_enforced_stats_w` is that limit reduced over this run's own 2 Hz trace and is the only honest cap cell; `power_limit_enforced_w` is the FIRST sample and is kept only so the two can be compared. A single wattage in a cap column for this posture is a fabrication.",
                  "cap_cell": "150 W, flat",
                  "cap_cell_note": "the enforced limit read 150 W on every one of 3 samples of this run's 2 Hz trace. Under dynamic-boost that is a FINDING, not a range.",
                  "memory_used_mib": {
                    "min": 19176.0,
                    "median": 19176.0,
                    "max": 19176.0,
                    "mean": 19176.0,
                    "n": 3
                  },
                  "utilization_pct": {
                    "min": 93.0,
                    "median": 93.0,
                    "max": 93.0,
                    "mean": 93.0,
                    "n": 3
                  },
                  "temperature_c": {
                    "min": 44.0,
                    "median": 44.0,
                    "max": 45.0,
                    "mean": 44.3,
                    "n": 3
                  },
                  "memory_temperature_c": null,
                  "fan_pct": null,
                  "pcie_width_current": {
                    "min": 8.0,
                    "median": 8.0,
                    "max": 8.0,
                    "mean": 8.0,
                    "n": 3
                  },
                  "clocks_sm_mhz": {
                    "min": 2220.0,
                    "median": 2250.0,
                    "max": 2340.0,
                    "mean": 2270.0,
                    "n": 3
                  },
                  "clocks_mem_mhz": {
                    "min": 14001.0,
                    "median": 14001.0,
                    "max": 14001.0,
                    "mean": 14001.0,
                    "n": 3
                  },
                  "throttle_sw_power_cap_fraction": 1.0,
                  "throttle_hw_slowdown_fraction": 0.0,
                  "throttle_sw_thermal_fraction": 0.0,
                  "throttle_hw_thermal_fraction": 0.0,
                  "n": 3
                },
                "busy_samples": 3
              }
            }
          },
          "energy_j_per_1k_tokens": 1136.29,
          "thermal": {
            "rule_version": 2,
            "measured_over": "samples where the GPU was busy (utilisation > 0)",
            "busy_samples": 3,
            "temp_max_c": 45.0,
            "temp_min_c": 44.0,
            "temp_rise_c": 1.0,
            "memory_temp_max_c": null,
            "memory_temp_median_c": null,
            "memory_temp_readable": false,
            "pcie_width_min": 8.0,
            "pcie_width_max": 8.0,
            "pcie_width_median": 8.0,
            "fan_mean_pct": null,
            "fan_max_pct": null,
            "clock_floor_mhz": 2220.0,
            "clock_max_mhz": 2340.0,
            "clock_median_mhz": 2250.0,
            "mem_clock_floor_mhz": 14001.0,
            "mem_clock_median_mhz": 14001.0,
            "mem_clock_max_mhz": 14001.0,
            "mem_clock_held": true,
            "memory_bus_width_bits": 256,
            "memory_technology": "GDDR7",
            "memory_bits_per_clock": null,
            "peak_mem_bandwidth_gbs_at_median_clock": null,
            "peak_mem_bandwidth_gbs_at_floor_clock": null,
            "peak_mem_bandwidth_gbs_at_max_clock": null,
            "peak_mem_bandwidth_withheld_because": "this harness's peak-bandwidth formula assumes two bits per memory clock per pin, which is exact for GDDR6 and GDDR6X and is checked against three published bandwidths. It has no receipt for the bits-per-clock of this board's memory, so the DERIVED bytes-per-second ceiling is withheld rather than computed from an unverified constant. The memory clock itself is a reading and is published unchanged.",
            "peak_mem_bandwidth_formula": "GB/s = clocks.mem MHz x 1e6 x <bits-per-clock for THIS board's memory> x (bus bits / 8) / 1e9; a CEILING at the clock that was observed, never an achieved rate. The bits-per-clock is registered per board and is NOT defaulted.",
            "clock_dropped": true,
            "power_limit_w": 150.0,
            "power_mean_w": 146.96,
            "power_pct_of_limit": 98.0,
            "power_pinned_at_limit": true,
            "sw_power_cap_active_fraction": 1.0,
            "hw_slowdown_active_fraction": 0.0,
            "sw_thermal_active_fraction": 0.0,
            "hw_thermal_active_fraction": 0.0,
            "at_throttling_temperature": false,
            "stop_core_temp_c": 83.0,
            "stop_memory_temp_c": 100.0,
            "hit_core_temp_stop": false,
            "hit_memory_temp_stop": false,
            "throttle_temperature_c": 80.0,
            "temp_rising_within_run": false,
            "verdict": "the POWER CAP doing its job: the SM clock fell while draw sat at 98.0% of the 150 W cap, at 45 C -- well below any throttling temperature"
          }
        },
        {
          "level": 1,
          "wall_s": 2.0446,
          "streams": [
            {
              "ttft_ms": 396.02,
              "first_token_was_thinking": false,
              "wall_s": 2.0441,
              "eval_count": 256,
              "eval_duration_ns": 1645107000,
              "decode_tok_s": 155.613,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 71317000,
              "prefill_tok_s": 7515.74,
              "load_duration_ns": 319053481,
              "total_duration_ns": 2040291133,
              "streamed_chunks": 256,
              "response_chars": 1044,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 154.728,
              "decode_tok_s_wall_gross": 125.24,
              "num_ctx_option": 4096,
              "stream": 1
            }
          ],
          "streams_ok": 1,
          "streams_failed": 0,
          "stream_errors": [],
          "trial_void": false,
          "per_stream_decode_tok_s": [
            155.613
          ],
          "per_stream_ttft_ms": [
            396.02
          ],
          "aggregate_sum_tok_s": 155.613,
          "aggregate_wall_tok_s": 125.207,
          "tokens_decoded_total": 256,
          "per_stream_decode": {
            "min": 155.613,
            "median": 155.613,
            "max": 155.613,
            "mean": 155.613,
            "n": 1
          },
          "per_stream_ttft": {
            "min": 396.02,
            "median": 396.02,
            "max": 396.02,
            "mean": 396.02,
            "n": 1
          },
          "power": {
            "idle_baseline": {
              "seconds": 3.0,
              "samples": 6,
              "power_w": {
                "min": 22.69,
                "median": 40.25,
                "max": 146.04,
                "mean": 60.7,
                "n": 6
              },
              "utilization_pct": {
                "min": 0.0,
                "median": 0.0,
                "max": 93.0,
                "mean": 15.5,
                "n": 6
              },
              "memory_used_mib": {
                "min": 19176.0,
                "median": 19176.0,
                "max": 19176.0,
                "mean": 19176.0,
                "n": 6
              },
              "power_w_per_card": {
                "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": 60.7
              },
              "memory_used_mib_per_card": {
                "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": 19176.0
              },
              "power_w_all_cards": 60.7,
              "ups_whole_box": {
                "load_pct": null,
                "realpower_w": null,
                "realpower_nominal_w": null,
                "status": null,
                "note": "the UPS's reading of EVERYTHING it feeds -- the whole box, not the cards. The card figures beside it are board power from nvidia-smi. Neither is the other."
              }
            },
            "mean_w": 118.63,
            "max_w": 146.58,
            "over_idle_w": 57.93,
            "samples": 8,
            "per_gpu": {
              "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": {
                "index": "0",
                "name": "NVIDIA GeForce RTX 5090 Laptop GPU",
                "power_w": {
                  "min": 22.71,
                  "median": 78.19,
                  "max": 146.58,
                  "mean": 84.39,
                  "n": 8
                },
                "power_limit_w": 150.0,
                "power_limit_field": "enforced.power.limit",
                "power_limit_enforced_w": 150.0,
                "cap": "dynamic-boost",
                "power_posture": "dynamic-boost",
                "power_limit_enforced_stats_w": {
                  "min": 150.0,
                  "median": 150.0,
                  "max": 150.0,
                  "mean": 150.0,
                  "n": 8
                },
                "power_posture_note": "dynamic-boost -- nvidia-powerd running and no cap set: the board floats between its power.default_limit and its power.max_limit against the CPU's draw, DURING the scored run. `power_limit_enforced_stats_w` is that limit reduced over this run's own 2 Hz trace and is the only honest cap cell; `power_limit_enforced_w` is the FIRST sample and is kept only so the two can be compared. A single wattage in a cap column for this posture is a fabrication.",
                "cap_cell": "150 W, flat",
                "cap_cell_note": "the enforced limit read 150 W on every one of 8 samples of this run's 2 Hz trace. Under dynamic-boost that is a FINDING, not a range.",
                "memory_used_mib": {
                  "min": 19176.0,
                  "median": 19176.0,
                  "max": 19176.0,
                  "mean": 19176.0,
                  "n": 8
                },
                "utilization_pct": {
                  "min": 0.0,
                  "median": 92.0,
                  "max": 93.0,
                  "mean": 57.9,
                  "n": 8
                },
                "temperature_c": {
                  "min": 37.0,
                  "median": 43.5,
                  "max": 45.0,
                  "mean": 42.1,
                  "n": 8
                },
                "memory_temperature_c": null,
                "fan_pct": null,
                "pcie_width_current": {
                  "min": 8.0,
                  "median": 8.0,
                  "max": 8.0,
                  "mean": 8.0,
                  "n": 8
                },
                "clocks_sm_mhz": {
                  "min": 1590.0,
                  "median": 2231.0,
                  "max": 2362.0,
                  "mean": 2108.1,
                  "n": 8
                },
                "clocks_mem_mhz": {
                  "min": 14001.0,
                  "median": 14001.0,
                  "max": 14001.0,
                  "mean": 14001.0,
                  "n": 8
                },
                "throttle_sw_power_cap_fraction": 0.625,
                "throttle_hw_slowdown_fraction": 0.0,
                "throttle_sw_thermal_fraction": 0.0,
                "throttle_hw_thermal_fraction": 0.0,
                "n": 8,
                "busy": {
                  "index": "0",
                  "name": "NVIDIA GeForce RTX 5090 Laptop GPU",
                  "power_w": {
                    "min": 54.78,
                    "median": 143.68,
                    "max": 146.58,
                    "mean": 118.63,
                    "n": 5
                  },
                  "power_limit_w": 150.0,
                  "power_limit_field": "enforced.power.limit",
                  "power_limit_enforced_w": 150.0,
                  "cap": "dynamic-boost",
                  "power_posture": "dynamic-boost",
                  "power_limit_enforced_stats_w": {
                    "min": 150.0,
                    "median": 150.0,
                    "max": 150.0,
                    "mean": 150.0,
                    "n": 5
                  },
                  "power_posture_note": "dynamic-boost -- nvidia-powerd running and no cap set: the board floats between its power.default_limit and its power.max_limit against the CPU's draw, DURING the scored run. `power_limit_enforced_stats_w` is that limit reduced over this run's own 2 Hz trace and is the only honest cap cell; `power_limit_enforced_w` is the FIRST sample and is kept only so the two can be compared. A single wattage in a cap column for this posture is a fabrication.",
                  "cap_cell": "150 W, flat",
                  "cap_cell_note": "the enforced limit read 150 W on every one of 5 samples of this run's 2 Hz trace. Under dynamic-boost that is a FINDING, not a range.",
                  "memory_used_mib": {
                    "min": 19176.0,
                    "median": 19176.0,
                    "max": 19176.0,
                    "mean": 19176.0,
                    "n": 5
                  },
                  "utilization_pct": {
                    "min": 92.0,
                    "median": 93.0,
                    "max": 93.0,
                    "mean": 92.6,
                    "n": 5
                  },
                  "temperature_c": {
                    "min": 43.0,
                    "median": 44.0,
                    "max": 45.0,
                    "mean": 44.2,
                    "n": 5
                  },
                  "memory_temperature_c": null,
                  "fan_pct": null,
                  "pcie_width_current": {
                    "min": 8.0,
                    "median": 8.0,
                    "max": 8.0,
                    "mean": 8.0,
                    "n": 5
                  },
                  "clocks_sm_mhz": {
                    "min": 2205.0,
                    "median": 2332.0,
                    "max": 2362.0,
                    "mean": 2303.6,
                    "n": 5
                  },
                  "clocks_mem_mhz": {
                    "min": 14001.0,
                    "median": 14001.0,
                    "max": 14001.0,
                    "mean": 14001.0,
                    "n": 5
                  },
                  "throttle_sw_power_cap_fraction": 1.0,
                  "throttle_hw_slowdown_fraction": 0.0,
                  "throttle_sw_thermal_fraction": 0.0,
                  "throttle_hw_thermal_fraction": 0.0,
                  "n": 5
                },
                "busy_samples": 5
              }
            }
          },
          "energy_j_per_1k_tokens": 947.47,
          "thermal": {
            "rule_version": 2,
            "measured_over": "samples where the GPU was busy (utilisation > 0)",
            "busy_samples": 5,
            "temp_max_c": 45.0,
            "temp_min_c": 43.0,
            "temp_rise_c": 2.0,
            "memory_temp_max_c": null,
            "memory_temp_median_c": null,
            "memory_temp_readable": false,
            "pcie_width_min": 8.0,
            "pcie_width_max": 8.0,
            "pcie_width_median": 8.0,
            "fan_mean_pct": null,
            "fan_max_pct": null,
            "clock_floor_mhz": 2205.0,
            "clock_max_mhz": 2362.0,
            "clock_median_mhz": 2332.0,
            "mem_clock_floor_mhz": 14001.0,
            "mem_clock_median_mhz": 14001.0,
            "mem_clock_max_mhz": 14001.0,
            "mem_clock_held": true,
            "memory_bus_width_bits": 256,
            "memory_technology": "GDDR7",
            "memory_bits_per_clock": null,
            "peak_mem_bandwidth_gbs_at_median_clock": null,
            "peak_mem_bandwidth_gbs_at_floor_clock": null,
            "peak_mem_bandwidth_gbs_at_max_clock": null,
            "peak_mem_bandwidth_withheld_because": "this harness's peak-bandwidth formula assumes two bits per memory clock per pin, which is exact for GDDR6 and GDDR6X and is checked against three published bandwidths. It has no receipt for the bits-per-clock of this board's memory, so the DERIVED bytes-per-second ceiling is withheld rather than computed from an unverified constant. The memory clock itself is a reading and is published unchanged.",
            "peak_mem_bandwidth_formula": "GB/s = clocks.mem MHz x 1e6 x <bits-per-clock for THIS board's memory> x (bus bits / 8) / 1e9; a CEILING at the clock that was observed, never an achieved rate. The bits-per-clock is registered per board and is NOT defaulted.",
            "clock_dropped": true,
            "power_limit_w": 150.0,
            "power_mean_w": 118.63,
            "power_pct_of_limit": 79.1,
            "power_pinned_at_limit": false,
            "sw_power_cap_active_fraction": 1.0,
            "hw_slowdown_active_fraction": 0.0,
            "sw_thermal_active_fraction": 0.0,
            "hw_thermal_active_fraction": 0.0,
            "at_throttling_temperature": false,
            "stop_core_temp_c": 83.0,
            "stop_memory_temp_c": 100.0,
            "hit_core_temp_stop": false,
            "hit_memory_temp_stop": false,
            "throttle_temperature_c": 80.0,
            "temp_rising_within_run": false,
            "verdict": "NEITHER cap: the SM clock varied with draw at 79.1% of the cap and the card at 45 C. On a memory-bound decode the clock follows the work, and nothing here was limiting it"
          }
        }
      ],
      "spread_gate": {
        "limit": 0.15,
        "values": [
          155.548,
          155.605,
          155.613
        ],
        "spread_fraction": 0.0004,
        "verdict": "spread 0.0% of the median within the 15% gate"
      },
      "voided_by_gate": false,
      "derived": {
        "aggregate_sum_tok_s": {
          "min": 155.548,
          "median": 155.605,
          "max": 155.613,
          "mean": 155.589,
          "n": 3
        },
        "aggregate_wall_tok_s": {
          "min": 125.207,
          "median": 125.73,
          "max": 129.333,
          "mean": 126.757,
          "n": 3
        },
        "per_stream_decode_tok_s": {
          "min": 155.548,
          "median": 155.605,
          "max": 155.613,
          "mean": 155.589,
          "n": 3
        },
        "per_stream_ttft_ms": {
          "min": 329.71,
          "median": 386.95,
          "max": 396.02,
          "mean": 370.89,
          "n": 3
        },
        "power_mean_w": {
          "min": 118.63,
          "median": 146.96,
          "max": 148.9,
          "mean": 138.16,
          "n": 3
        },
        "energy_j_per_1k_tokens": {
          "min": 947.47,
          "median": 1136.29,
          "max": 1184.28,
          "mean": 1089.35,
          "n": 3
        },
        "temp_max_c": {
          "min": 44.0,
          "median": 45.0,
          "max": 45.0,
          "mean": 44.7,
          "n": 3
        },
        "fan_mean_pct": null,
        "clock_floor_mhz": {
          "min": 2205.0,
          "median": 2220.0,
          "max": 2242.0,
          "mean": 2222.3,
          "n": 3
        }
      },
      "thermal_verdict": "the POWER CAP doing its job: the SM clock fell while draw sat at 106.4% of the 140 W cap, at 44 C -- well below any throttling temperature; the POWER CAP doing its job: the SM clock fell while draw sat at 98.0% of the 150 W cap, at 45 C -- well below any throttling temperature; NEITHER cap: the SM clock varied with draw at 79.1% of the cap and the card at 45 C. On a memory-bound decode the clock follows the work, and nothing here was limiting it"
    },
    {
      "level": 2,
      "contention_gate": {
        "attempts": [
          {
            "window_s": 10.0,
            "samples": 20,
            "cards": {
              "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": {
                "utilization_pct": {
                  "min": 0.0,
                  "median": 0.0,
                  "max": 93.0,
                  "mean": 4.7,
                  "n": 20
                },
                "power_w": {
                  "min": 22.65,
                  "median": 22.7,
                  "max": 145.36,
                  "mean": 33.57,
                  "n": 20
                },
                "card_seen": true
              }
            },
            "max_mean_util_pct": 4.7,
            "all_cards_seen": true,
            "attempt": 1,
            "verdict": "quiet: busiest card mean 4.7% <= 5.0%"
          }
        ],
        "passed": true,
        "cards": [
          "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff"
        ]
      },
      "trials": [
        {
          "level": 2,
          "wall_s": 2.5194,
          "streams": [
            {
              "ttft_ms": 425.2,
              "first_token_was_thinking": false,
              "wall_s": 2.5006,
              "eval_count": 256,
              "eval_duration_ns": 2072792000,
              "decode_tok_s": 123.505,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 105842000,
              "prefill_tok_s": 5064.152,
              "load_duration_ns": 313324354,
              "total_duration_ns": 2497076810,
              "streamed_chunks": 256,
              "response_chars": 1026,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 122.867,
              "decode_tok_s_wall_gross": 102.374,
              "num_ctx_option": 4096,
              "stream": 1
            },
            {
              "ttft_ms": 569.84,
              "first_token_was_thinking": false,
              "wall_s": 2.5189,
              "eval_count": 256,
              "eval_duration_ns": 1946516000,
              "decode_tok_s": 131.517,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 211378000,
              "prefill_tok_s": 2535.742,
              "load_duration_ns": 338657223,
              "total_duration_ns": 2515572332,
              "streamed_chunks": 256,
              "response_chars": 1041,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 130.833,
              "decode_tok_s_wall_gross": 101.632,
              "num_ctx_option": 4096,
              "stream": 2
            }
          ],
          "streams_ok": 2,
          "streams_failed": 0,
          "stream_errors": [],
          "trial_void": false,
          "per_stream_decode_tok_s": [
            123.505,
            131.517
          ],
          "per_stream_ttft_ms": [
            425.2,
            569.84
          ],
          "aggregate_sum_tok_s": 255.022,
          "aggregate_wall_tok_s": 203.219,
          "tokens_decoded_total": 512,
          "per_stream_decode": {
            "min": 123.505,
            "median": 127.511,
            "max": 131.517,
            "mean": 127.511,
            "n": 2
          },
          "per_stream_ttft": {
            "min": 425.2,
            "median": 497.52,
            "max": 569.84,
            "mean": 497.52,
            "n": 2
          },
          "power": {
            "idle_baseline": {
              "seconds": 3.0,
              "samples": 6,
              "power_w": {
                "min": 22.65,
                "median": 22.66,
                "max": 22.68,
                "mean": 22.66,
                "n": 6
              },
              "utilization_pct": {
                "min": 0.0,
                "median": 0.0,
                "max": 0.0,
                "mean": 0.0,
                "n": 6
              },
              "memory_used_mib": {
                "min": 19176.0,
                "median": 19176.0,
                "max": 19176.0,
                "mean": 19176.0,
                "n": 6
              },
              "power_w_per_card": {
                "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": 22.66
              },
              "memory_used_mib_per_card": {
                "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": 19176.0
              },
              "power_w_all_cards": 22.66,
              "ups_whole_box": {
                "load_pct": null,
                "realpower_w": null,
                "realpower_nominal_w": null,
                "status": null,
                "note": "the UPS's reading of EVERYTHING it feeds -- the whole box, not the cards. The card figures beside it are board power from nvidia-smi. Neither is the other."
              }
            },
            "mean_w": 144.35,
            "max_w": 146.0,
            "over_idle_w": 121.69,
            "samples": 10,
            "per_gpu": {
              "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": {
                "index": "0",
                "name": "NVIDIA GeForce RTX 5090 Laptop GPU",
                "power_w": {
                  "min": 22.66,
                  "median": 118.43,
                  "max": 146.0,
                  "mean": 94.69,
                  "n": 10
                },
                "power_limit_w": 150.0,
                "power_limit_field": "enforced.power.limit",
                "power_limit_enforced_w": 150.0,
                "cap": "dynamic-boost",
                "power_posture": "dynamic-boost",
                "power_limit_enforced_stats_w": {
                  "min": 150.0,
                  "median": 150.0,
                  "max": 150.0,
                  "mean": 150.0,
                  "n": 10
                },
                "power_posture_note": "dynamic-boost -- nvidia-powerd running and no cap set: the board floats between its power.default_limit and its power.max_limit against the CPU's draw, DURING the scored run. `power_limit_enforced_stats_w` is that limit reduced over this run's own 2 Hz trace and is the only honest cap cell; `power_limit_enforced_w` is the FIRST sample and is kept only so the two can be compared. A single wattage in a cap column for this posture is a fabrication.",
                "cap_cell": "150 W, flat",
                "cap_cell_note": "the enforced limit read 150 W on every one of 10 samples of this run's 2 Hz trace. Under dynamic-boost that is a FINDING, not a range.",
                "memory_used_mib": {
                  "min": 19176.0,
                  "median": 19176.0,
                  "max": 19176.0,
                  "mean": 19176.0,
                  "n": 10
                },
                "utilization_pct": {
                  "min": 0.0,
                  "median": 44.0,
                  "max": 89.0,
                  "mean": 44.3,
                  "n": 10
                },
                "temperature_c": {
                  "min": 36.0,
                  "median": 44.0,
                  "max": 45.0,
                  "mean": 42.5,
                  "n": 10
                },
                "memory_temperature_c": null,
                "fan_pct": null,
                "pcie_width_current": {
                  "min": 8.0,
                  "median": 8.0,
                  "max": 8.0,
                  "mean": 8.0,
                  "n": 10
                },
                "clocks_sm_mhz": {
                  "min": 1590.0,
                  "median": 2205.0,
                  "max": 2265.0,
                  "mean": 2063.2,
                  "n": 10
                },
                "clocks_mem_mhz": {
                  "min": 14001.0,
                  "median": 14001.0,
                  "max": 14001.0,
                  "mean": 14001.0,
                  "n": 10
                },
                "throttle_sw_power_cap_fraction": 0.7,
                "throttle_hw_slowdown_fraction": 0.0,
                "throttle_sw_thermal_fraction": 0.0,
                "throttle_hw_thermal_fraction": 0.0,
                "n": 10,
                "busy": {
                  "index": "0",
                  "name": "NVIDIA GeForce RTX 5090 Laptop GPU",
                  "power_w": {
                    "min": 139.7,
                    "median": 145.84,
                    "max": 146.0,
                    "mean": 144.35,
                    "n": 5
                  },
                  "power_limit_w": 150.0,
                  "power_limit_field": "enforced.power.limit",
                  "power_limit_enforced_w": 150.0,
                  "cap": "dynamic-boost",
                  "power_posture": "dynamic-boost",
                  "power_limit_enforced_stats_w": {
                    "min": 150.0,
                    "median": 150.0,
                    "max": 150.0,
                    "mean": 150.0,
                    "n": 5
                  },
                  "power_posture_note": "dynamic-boost -- nvidia-powerd running and no cap set: the board floats between its power.default_limit and its power.max_limit against the CPU's draw, DURING the scored run. `power_limit_enforced_stats_w` is that limit reduced over this run's own 2 Hz trace and is the only honest cap cell; `power_limit_enforced_w` is the FIRST sample and is kept only so the two can be compared. A single wattage in a cap column for this posture is a fabrication.",
                  "cap_cell": "150 W, flat",
                  "cap_cell_note": "the enforced limit read 150 W on every one of 5 samples of this run's 2 Hz trace. Under dynamic-boost that is a FINDING, not a range.",
                  "memory_used_mib": {
                    "min": 19176.0,
                    "median": 19176.0,
                    "max": 19176.0,
                    "mean": 19176.0,
                    "n": 5
                  },
                  "utilization_pct": {
                    "min": 88.0,
                    "median": 89.0,
                    "max": 89.0,
                    "mean": 88.6,
                    "n": 5
                  },
                  "temperature_c": {
                    "min": 44.0,
                    "median": 45.0,
                    "max": 45.0,
                    "mean": 44.6,
                    "n": 5
                  },
                  "memory_temperature_c": null,
                  "fan_pct": null,
                  "pcie_width_current": {
                    "min": 8.0,
                    "median": 8.0,
                    "max": 8.0,
                    "mean": 8.0,
                    "n": 5
                  },
                  "clocks_sm_mhz": {
                    "min": 2190.0,
                    "median": 2235.0,
                    "max": 2242.0,
                    "mean": 2221.4,
                    "n": 5
                  },
                  "clocks_mem_mhz": {
                    "min": 14001.0,
                    "median": 14001.0,
                    "max": 14001.0,
                    "mean": 14001.0,
                    "n": 5
                  },
                  "throttle_sw_power_cap_fraction": 1.0,
                  "throttle_hw_slowdown_fraction": 0.0,
                  "throttle_sw_thermal_fraction": 0.0,
                  "throttle_hw_thermal_fraction": 0.0,
                  "n": 5
                },
                "busy_samples": 5
              }
            }
          },
          "energy_j_per_1k_tokens": 710.32,
          "thermal": {
            "rule_version": 2,
            "measured_over": "samples where the GPU was busy (utilisation > 0)",
            "busy_samples": 5,
            "temp_max_c": 45.0,
            "temp_min_c": 44.0,
            "temp_rise_c": 1.0,
            "memory_temp_max_c": null,
            "memory_temp_median_c": null,
            "memory_temp_readable": false,
            "pcie_width_min": 8.0,
            "pcie_width_max": 8.0,
            "pcie_width_median": 8.0,
            "fan_mean_pct": null,
            "fan_max_pct": null,
            "clock_floor_mhz": 2190.0,
            "clock_max_mhz": 2242.0,
            "clock_median_mhz": 2235.0,
            "mem_clock_floor_mhz": 14001.0,
            "mem_clock_median_mhz": 14001.0,
            "mem_clock_max_mhz": 14001.0,
            "mem_clock_held": true,
            "memory_bus_width_bits": 256,
            "memory_technology": "GDDR7",
            "memory_bits_per_clock": null,
            "peak_mem_bandwidth_gbs_at_median_clock": null,
            "peak_mem_bandwidth_gbs_at_floor_clock": null,
            "peak_mem_bandwidth_gbs_at_max_clock": null,
            "peak_mem_bandwidth_withheld_because": "this harness's peak-bandwidth formula assumes two bits per memory clock per pin, which is exact for GDDR6 and GDDR6X and is checked against three published bandwidths. It has no receipt for the bits-per-clock of this board's memory, so the DERIVED bytes-per-second ceiling is withheld rather than computed from an unverified constant. The memory clock itself is a reading and is published unchanged.",
            "peak_mem_bandwidth_formula": "GB/s = clocks.mem MHz x 1e6 x <bits-per-clock for THIS board's memory> x (bus bits / 8) / 1e9; a CEILING at the clock that was observed, never an achieved rate. The bits-per-clock is registered per board and is NOT defaulted.",
            "clock_dropped": false,
            "power_limit_w": 150.0,
            "power_mean_w": 144.35,
            "power_pct_of_limit": 96.2,
            "power_pinned_at_limit": true,
            "sw_power_cap_active_fraction": 1.0,
            "hw_slowdown_active_fraction": 0.0,
            "sw_thermal_active_fraction": 0.0,
            "hw_thermal_active_fraction": 0.0,
            "at_throttling_temperature": false,
            "stop_core_temp_c": 83.0,
            "stop_memory_temp_c": 100.0,
            "hit_core_temp_stop": false,
            "hit_memory_temp_stop": false,
            "throttle_temperature_c": 80.0,
            "temp_rising_within_run": false,
            "verdict": "no clock drop during decode (draw held at 96.2% of the 150 W cap)"
          }
        },
        {
          "level": 2,
          "wall_s": 2.4775,
          "streams": [
            {
              "ttft_ms": 493.09,
              "first_token_was_thinking": false,
              "wall_s": 2.4769,
              "eval_count": 256,
              "eval_duration_ns": 1980747000,
              "decode_tok_s": 129.244,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 121985000,
              "prefill_tok_s": 4393.983,
              "load_duration_ns": 293718771,
              "total_duration_ns": 2473242117,
              "streamed_chunks": 256,
              "response_chars": 1040,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 128.539,
              "decode_tok_s_wall_gross": 103.354,
              "num_ctx_option": 4096,
              "stream": 1
            },
            {
              "ttft_ms": 494.11,
              "first_token_was_thinking": false,
              "wall_s": 2.4765,
              "eval_count": 256,
              "eval_duration_ns": 1979827000,
              "decode_tok_s": 129.304,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 94073000,
              "prefill_tok_s": 5697.703,
              "load_duration_ns": 312226399,
              "total_duration_ns": 2473340001,
              "streamed_chunks": 256,
              "response_chars": 1031,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 128.634,
              "decode_tok_s_wall_gross": 103.373,
              "num_ctx_option": 4096,
              "stream": 2
            }
          ],
          "streams_ok": 2,
          "streams_failed": 0,
          "stream_errors": [],
          "trial_void": false,
          "per_stream_decode_tok_s": [
            129.244,
            129.304
          ],
          "per_stream_ttft_ms": [
            493.09,
            494.11
          ],
          "aggregate_sum_tok_s": 258.548,
          "aggregate_wall_tok_s": 206.658,
          "tokens_decoded_total": 512,
          "per_stream_decode": {
            "min": 129.244,
            "median": 129.274,
            "max": 129.304,
            "mean": 129.274,
            "n": 2
          },
          "per_stream_ttft": {
            "min": 493.09,
            "median": 493.6,
            "max": 494.11,
            "mean": 493.6,
            "n": 2
          },
          "power": {
            "idle_baseline": {
              "seconds": 3.0,
              "samples": 6,
              "power_w": {
                "min": 22.82,
                "median": 35.6,
                "max": 145.91,
                "mean": 58.15,
                "n": 6
              },
              "utilization_pct": {
                "min": 0.0,
                "median": 0.0,
                "max": 89.0,
                "mean": 14.8,
                "n": 6
              },
              "memory_used_mib": {
                "min": 19176.0,
                "median": 19176.0,
                "max": 19176.0,
                "mean": 19176.0,
                "n": 6
              },
              "power_w_per_card": {
                "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": 58.15
              },
              "memory_used_mib_per_card": {
                "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": 19176.0
              },
              "power_w_all_cards": 58.15,
              "ups_whole_box": {
                "load_pct": null,
                "realpower_w": null,
                "realpower_nominal_w": null,
                "status": null,
                "note": "the UPS's reading of EVERYTHING it feeds -- the whole box, not the cards. The card figures beside it are board power from nvidia-smi. Neither is the other."
              }
            },
            "mean_w": 118.57,
            "max_w": 146.91,
            "over_idle_w": 60.42,
            "samples": 9,
            "per_gpu": {
              "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": {
                "index": "0",
                "name": "NVIDIA GeForce RTX 5090 Laptop GPU",
                "power_w": {
                  "min": 22.81,
                  "median": 92.19,
                  "max": 146.91,
                  "mean": 86.73,
                  "n": 9
                },
                "power_limit_w": 150.0,
                "power_limit_field": "enforced.power.limit",
                "power_limit_enforced_w": 150.0,
                "cap": "dynamic-boost",
                "power_posture": "dynamic-boost",
                "power_limit_enforced_stats_w": {
                  "min": 150.0,
                  "median": 150.0,
                  "max": 150.0,
                  "mean": 150.0,
                  "n": 9
                },
                "power_posture_note": "dynamic-boost -- nvidia-powerd running and no cap set: the board floats between its power.default_limit and its power.max_limit against the CPU's draw, DURING the scored run. `power_limit_enforced_stats_w` is that limit reduced over this run's own 2 Hz trace and is the only honest cap cell; `power_limit_enforced_w` is the FIRST sample and is kept only so the two can be compared. A single wattage in a cap column for this posture is a fabrication.",
                "cap_cell": "150 W, flat",
                "cap_cell_note": "the enforced limit read 150 W on every one of 9 samples of this run's 2 Hz trace. Under dynamic-boost that is a FINDING, not a range.",
                "memory_used_mib": {
                  "min": 19176.0,
                  "median": 19176.0,
                  "max": 19176.0,
                  "mean": 19176.0,
                  "n": 9
                },
                "utilization_pct": {
                  "min": 0.0,
                  "median": 89.0,
                  "max": 90.0,
                  "mean": 59.8,
                  "n": 9
                },
                "temperature_c": {
                  "min": 37.0,
                  "median": 44.0,
                  "max": 45.0,
                  "mean": 42.0,
                  "n": 9
                },
                "memory_temperature_c": null,
                "fan_pct": null,
                "pcie_width_current": {
                  "min": 8.0,
                  "median": 8.0,
                  "max": 8.0,
                  "mean": 8.0,
                  "n": 9
                },
                "clocks_sm_mhz": {
                  "min": 1590.0,
                  "median": 2205.0,
                  "max": 2265.0,
                  "mean": 2027.3,
                  "n": 9
                },
                "clocks_mem_mhz": {
                  "min": 14001.0,
                  "median": 14001.0,
                  "max": 14001.0,
                  "mean": 14001.0,
                  "n": 9
                },
                "throttle_sw_power_cap_fraction": 0.6667,
                "throttle_hw_slowdown_fraction": 0.0,
                "throttle_sw_thermal_fraction": 0.0,
                "throttle_hw_thermal_fraction": 0.0,
                "n": 9,
                "busy": {
                  "index": "0",
                  "name": "NVIDIA GeForce RTX 5090 Laptop GPU",
                  "power_w": {
                    "min": 44.29,
                    "median": 140.89,
                    "max": 146.91,
                    "mean": 118.57,
                    "n": 6
                  },
                  "power_limit_w": 150.0,
                  "power_limit_field": "enforced.power.limit",
                  "power_limit_enforced_w": 150.0,
                  "cap": "dynamic-boost",
                  "power_posture": "dynamic-boost",
                  "power_limit_enforced_stats_w": {
                    "min": 150.0,
                    "median": 150.0,
                    "max": 150.0,
                    "mean": 150.0,
                    "n": 6
                  },
                  "power_posture_note": "dynamic-boost -- nvidia-powerd running and no cap set: the board floats between its power.default_limit and its power.max_limit against the CPU's draw, DURING the scored run. `power_limit_enforced_stats_w` is that limit reduced over this run's own 2 Hz trace and is the only honest cap cell; `power_limit_enforced_w` is the FIRST sample and is kept only so the two can be compared. A single wattage in a cap column for this posture is a fabrication.",
                  "cap_cell": "150 W, flat",
                  "cap_cell_note": "the enforced limit read 150 W on every one of 6 samples of this run's 2 Hz trace. Under dynamic-boost that is a FINDING, not a range.",
                  "memory_used_mib": {
                    "min": 19176.0,
                    "median": 19176.0,
                    "max": 19176.0,
                    "mean": 19176.0,
                    "n": 6
                  },
                  "utilization_pct": {
                    "min": 89.0,
                    "median": 90.0,
                    "max": 90.0,
                    "mean": 89.7,
                    "n": 6
                  },
                  "temperature_c": {
                    "min": 43.0,
                    "median": 45.0,
                    "max": 45.0,
                    "mean": 44.5,
                    "n": 6
                  },
                  "memory_temperature_c": null,
                  "fan_pct": null,
                  "pcie_width_current": {
                    "min": 8.0,
                    "median": 8.0,
                    "max": 8.0,
                    "mean": 8.0,
                    "n": 6
                  },
                  "clocks_sm_mhz": {
                    "min": 2197.0,
                    "median": 2242.5,
                    "max": 2265.0,
                    "mean": 2234.8,
                    "n": 6
                  },
                  "clocks_mem_mhz": {
                    "min": 14001.0,
                    "median": 14001.0,
                    "max": 14001.0,
                    "mean": 14001.0,
                    "n": 6
                  },
                  "throttle_sw_power_cap_fraction": 1.0,
                  "throttle_hw_slowdown_fraction": 0.0,
                  "throttle_sw_thermal_fraction": 0.0,
                  "throttle_hw_thermal_fraction": 0.0,
                  "n": 6
                },
                "busy_samples": 6
              }
            }
          },
          "energy_j_per_1k_tokens": 573.75,
          "thermal": {
            "rule_version": 2,
            "measured_over": "samples where the GPU was busy (utilisation > 0)",
            "busy_samples": 6,
            "temp_max_c": 45.0,
            "temp_min_c": 43.0,
            "temp_rise_c": 2.0,
            "memory_temp_max_c": null,
            "memory_temp_median_c": null,
            "memory_temp_readable": false,
            "pcie_width_min": 8.0,
            "pcie_width_max": 8.0,
            "pcie_width_median": 8.0,
            "fan_mean_pct": null,
            "fan_max_pct": null,
            "clock_floor_mhz": 2197.0,
            "clock_max_mhz": 2265.0,
            "clock_median_mhz": 2242.5,
            "mem_clock_floor_mhz": 14001.0,
            "mem_clock_median_mhz": 14001.0,
            "mem_clock_max_mhz": 14001.0,
            "mem_clock_held": true,
            "memory_bus_width_bits": 256,
            "memory_technology": "GDDR7",
            "memory_bits_per_clock": null,
            "peak_mem_bandwidth_gbs_at_median_clock": null,
            "peak_mem_bandwidth_gbs_at_floor_clock": null,
            "peak_mem_bandwidth_gbs_at_max_clock": null,
            "peak_mem_bandwidth_withheld_because": "this harness's peak-bandwidth formula assumes two bits per memory clock per pin, which is exact for GDDR6 and GDDR6X and is checked against three published bandwidths. It has no receipt for the bits-per-clock of this board's memory, so the DERIVED bytes-per-second ceiling is withheld rather than computed from an unverified constant. The memory clock itself is a reading and is published unchanged.",
            "peak_mem_bandwidth_formula": "GB/s = clocks.mem MHz x 1e6 x <bits-per-clock for THIS board's memory> x (bus bits / 8) / 1e9; a CEILING at the clock that was observed, never an achieved rate. The bits-per-clock is registered per board and is NOT defaulted.",
            "clock_dropped": false,
            "power_limit_w": 150.0,
            "power_mean_w": 118.57,
            "power_pct_of_limit": 79.0,
            "power_pinned_at_limit": false,
            "sw_power_cap_active_fraction": 1.0,
            "hw_slowdown_active_fraction": 0.0,
            "sw_thermal_active_fraction": 0.0,
            "hw_thermal_active_fraction": 0.0,
            "at_throttling_temperature": false,
            "stop_core_temp_c": 83.0,
            "stop_memory_temp_c": 100.0,
            "hit_core_temp_stop": false,
            "hit_memory_temp_stop": false,
            "throttle_temperature_c": 80.0,
            "temp_rising_within_run": false,
            "verdict": "no clock drop during decode"
          }
        },
        {
          "level": 2,
          "wall_s": 2.488,
          "streams": [
            {
              "ttft_ms": 506.49,
              "first_token_was_thinking": false,
              "wall_s": 2.4872,
              "eval_count": 256,
              "eval_duration_ns": 1978748000,
              "decode_tok_s": 129.375,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 96494000,
              "prefill_tok_s": 5554.75,
              "load_duration_ns": 339776068,
              "total_duration_ns": 2484516498,
              "streamed_chunks": 256,
              "response_chars": 1031,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 128.745,
              "decode_tok_s_wall_gross": 102.929,
              "num_ctx_option": 4096,
              "stream": 1
            },
            {
              "ttft_ms": 505.25,
              "first_token_was_thinking": false,
              "wall_s": 2.4878,
              "eval_count": 256,
              "eval_duration_ns": 1979746000,
              "decode_tok_s": 129.31,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 124950000,
              "prefill_tok_s": 4289.716,
              "load_duration_ns": 301912540,
              "total_duration_ns": 2484193056,
              "streamed_chunks": 256,
              "response_chars": 1040,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 128.625,
              "decode_tok_s_wall_gross": 102.904,
              "num_ctx_option": 4096,
              "stream": 2
            }
          ],
          "streams_ok": 2,
          "streams_failed": 0,
          "stream_errors": [],
          "trial_void": false,
          "per_stream_decode_tok_s": [
            129.375,
            129.31
          ],
          "per_stream_ttft_ms": [
            506.49,
            505.25
          ],
          "aggregate_sum_tok_s": 258.685,
          "aggregate_wall_tok_s": 205.79,
          "tokens_decoded_total": 512,
          "per_stream_decode": {
            "min": 129.31,
            "median": 129.343,
            "max": 129.375,
            "mean": 129.343,
            "n": 2
          },
          "per_stream_ttft": {
            "min": 505.25,
            "median": 505.87,
            "max": 506.49,
            "mean": 505.87,
            "n": 2
          },
          "power": {
            "idle_baseline": {
              "seconds": 3.0,
              "samples": 6,
              "power_w": {
                "min": 22.72,
                "median": 44.89,
                "max": 145.73,
                "mean": 61.88,
                "n": 6
              },
              "utilization_pct": {
                "min": 90.0,
                "median": 90.0,
                "max": 90.0,
                "mean": 90.0,
                "n": 6
              },
              "memory_used_mib": {
                "min": 19176.0,
                "median": 19176.0,
                "max": 19176.0,
                "mean": 19176.0,
                "n": 6
              },
              "power_w_per_card": {
                "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": 61.88
              },
              "memory_used_mib_per_card": {
                "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": 19176.0
              },
              "power_w_all_cards": 61.88,
              "ups_whole_box": {
                "load_pct": null,
                "realpower_w": null,
                "realpower_nominal_w": null,
                "status": null,
                "note": "the UPS's reading of EVERYTHING it feeds -- the whole box, not the cards. The card figures beside it are board power from nvidia-smi. Neither is the other."
              }
            },
            "mean_w": 115.88,
            "max_w": 146.76,
            "over_idle_w": 54.0,
            "samples": 10,
            "per_gpu": {
              "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": {
                "index": "0",
                "name": "NVIDIA GeForce RTX 5090 Laptop GPU",
                "power_w": {
                  "min": 22.71,
                  "median": 91.74,
                  "max": 146.76,
                  "mean": 88.08,
                  "n": 10
                },
                "power_limit_w": 150.0,
                "power_limit_field": "enforced.power.limit",
                "power_limit_enforced_w": 150.0,
                "cap": "dynamic-boost",
                "power_posture": "dynamic-boost",
                "power_limit_enforced_stats_w": {
                  "min": 150.0,
                  "median": 150.0,
                  "max": 150.0,
                  "mean": 150.0,
                  "n": 10
                },
                "power_posture_note": "dynamic-boost -- nvidia-powerd running and no cap set: the board floats between its power.default_limit and its power.max_limit against the CPU's draw, DURING the scored run. `power_limit_enforced_stats_w` is that limit reduced over this run's own 2 Hz trace and is the only honest cap cell; `power_limit_enforced_w` is the FIRST sample and is kept only so the two can be compared. A single wattage in a cap column for this posture is a fabrication.",
                "cap_cell": "150 W, flat",
                "cap_cell_note": "the enforced limit read 150 W on every one of 10 samples of this run's 2 Hz trace. Under dynamic-boost that is a FINDING, not a range.",
                "memory_used_mib": {
                  "min": 19176.0,
                  "median": 19176.0,
                  "max": 19176.0,
                  "mean": 19176.0,
                  "n": 10
                },
                "utilization_pct": {
                  "min": 0.0,
                  "median": 89.0,
                  "max": 90.0,
                  "mean": 62.7,
                  "n": 10
                },
                "temperature_c": {
                  "min": 38.0,
                  "median": 45.0,
                  "max": 46.0,
                  "mean": 43.2,
                  "n": 10
                },
                "memory_temperature_c": null,
                "fan_pct": null,
                "pcie_width_current": {
                  "min": 8.0,
                  "median": 8.0,
                  "max": 8.0,
                  "mean": 8.0,
                  "n": 10
                },
                "clocks_sm_mhz": {
                  "min": 1590.0,
                  "median": 2216.0,
                  "max": 2257.0,
                  "mean": 2045.8,
                  "n": 10
                },
                "clocks_mem_mhz": {
                  "min": 14001.0,
                  "median": 14001.0,
                  "max": 14001.0,
                  "mean": 14001.0,
                  "n": 10
                },
                "throttle_sw_power_cap_fraction": 0.7,
                "throttle_hw_slowdown_fraction": 0.0,
                "throttle_sw_thermal_fraction": 0.0,
                "throttle_hw_thermal_fraction": 0.0,
                "n": 10,
                "busy": {
                  "index": "0",
                  "name": "NVIDIA GeForce RTX 5090 Laptop GPU",
                  "power_w": {
                    "min": 43.62,
                    "median": 145.36,
                    "max": 146.76,
                    "mean": 115.88,
                    "n": 7
                  },
                  "power_limit_w": 150.0,
                  "power_limit_field": "enforced.power.limit",
                  "power_limit_enforced_w": 150.0,
                  "cap": "dynamic-boost",
                  "power_posture": "dynamic-boost",
                  "power_limit_enforced_stats_w": {
                    "min": 150.0,
                    "median": 150.0,
                    "max": 150.0,
                    "mean": 150.0,
                    "n": 7
                  },
                  "power_posture_note": "dynamic-boost -- nvidia-powerd running and no cap set: the board floats between its power.default_limit and its power.max_limit against the CPU's draw, DURING the scored run. `power_limit_enforced_stats_w` is that limit reduced over this run's own 2 Hz trace and is the only honest cap cell; `power_limit_enforced_w` is the FIRST sample and is kept only so the two can be compared. A single wattage in a cap column for this posture is a fabrication.",
                  "cap_cell": "150 W, flat",
                  "cap_cell_note": "the enforced limit read 150 W on every one of 7 samples of this run's 2 Hz trace. Under dynamic-boost that is a FINDING, not a range.",
                  "memory_used_mib": {
                    "min": 19176.0,
                    "median": 19176.0,
                    "max": 19176.0,
                    "mean": 19176.0,
                    "n": 7
                  },
                  "utilization_pct": {
                    "min": 89.0,
                    "median": 90.0,
                    "max": 90.0,
                    "mean": 89.6,
                    "n": 7
                  },
                  "temperature_c": {
                    "min": 44.0,
                    "median": 45.0,
                    "max": 46.0,
                    "mean": 45.1,
                    "n": 7
                  },
                  "memory_temperature_c": null,
                  "fan_pct": null,
                  "pcie_width_current": {
                    "min": 8.0,
                    "median": 8.0,
                    "max": 8.0,
                    "mean": 8.0,
                    "n": 7
                  },
                  "clocks_sm_mhz": {
                    "min": 2190.0,
                    "median": 2235.0,
                    "max": 2257.0,
                    "mean": 2228.3,
                    "n": 7
                  },
                  "clocks_mem_mhz": {
                    "min": 14001.0,
                    "median": 14001.0,
                    "max": 14001.0,
                    "mean": 14001.0,
                    "n": 7
                  },
                  "throttle_sw_power_cap_fraction": 1.0,
                  "throttle_hw_slowdown_fraction": 0.0,
                  "throttle_sw_thermal_fraction": 0.0,
                  "throttle_hw_thermal_fraction": 0.0,
                  "n": 7
                },
                "busy_samples": 7
              }
            }
          },
          "energy_j_per_1k_tokens": 563.1,
          "thermal": {
            "rule_version": 2,
            "measured_over": "samples where the GPU was busy (utilisation > 0)",
            "busy_samples": 7,
            "temp_max_c": 46.0,
            "temp_min_c": 44.0,
            "temp_rise_c": 2.0,
            "memory_temp_max_c": null,
            "memory_temp_median_c": null,
            "memory_temp_readable": false,
            "pcie_width_min": 8.0,
            "pcie_width_max": 8.0,
            "pcie_width_median": 8.0,
            "fan_mean_pct": null,
            "fan_max_pct": null,
            "clock_floor_mhz": 2190.0,
            "clock_max_mhz": 2257.0,
            "clock_median_mhz": 2235.0,
            "mem_clock_floor_mhz": 14001.0,
            "mem_clock_median_mhz": 14001.0,
            "mem_clock_max_mhz": 14001.0,
            "mem_clock_held": true,
            "memory_bus_width_bits": 256,
            "memory_technology": "GDDR7",
            "memory_bits_per_clock": null,
            "peak_mem_bandwidth_gbs_at_median_clock": null,
            "peak_mem_bandwidth_gbs_at_floor_clock": null,
            "peak_mem_bandwidth_gbs_at_max_clock": null,
            "peak_mem_bandwidth_withheld_because": "this harness's peak-bandwidth formula assumes two bits per memory clock per pin, which is exact for GDDR6 and GDDR6X and is checked against three published bandwidths. It has no receipt for the bits-per-clock of this board's memory, so the DERIVED bytes-per-second ceiling is withheld rather than computed from an unverified constant. The memory clock itself is a reading and is published unchanged.",
            "peak_mem_bandwidth_formula": "GB/s = clocks.mem MHz x 1e6 x <bits-per-clock for THIS board's memory> x (bus bits / 8) / 1e9; a CEILING at the clock that was observed, never an achieved rate. The bits-per-clock is registered per board and is NOT defaulted.",
            "clock_dropped": false,
            "power_limit_w": 150.0,
            "power_mean_w": 115.88,
            "power_pct_of_limit": 77.3,
            "power_pinned_at_limit": false,
            "sw_power_cap_active_fraction": 1.0,
            "hw_slowdown_active_fraction": 0.0,
            "sw_thermal_active_fraction": 0.0,
            "hw_thermal_active_fraction": 0.0,
            "at_throttling_temperature": false,
            "stop_core_temp_c": 83.0,
            "stop_memory_temp_c": 100.0,
            "hit_core_temp_stop": false,
            "hit_memory_temp_stop": false,
            "throttle_temperature_c": 80.0,
            "temp_rising_within_run": false,
            "verdict": "no clock drop during decode"
          }
        }
      ],
      "spread_gate": {
        "limit": 0.15,
        "values": [
          255.022,
          258.548,
          258.685
        ],
        "spread_fraction": 0.0142,
        "verdict": "spread 1.4% of the median within the 15% gate"
      },
      "voided_by_gate": false,
      "derived": {
        "aggregate_sum_tok_s": {
          "min": 255.022,
          "median": 258.548,
          "max": 258.685,
          "mean": 257.418,
          "n": 3
        },
        "aggregate_wall_tok_s": {
          "min": 203.219,
          "median": 205.79,
          "max": 206.658,
          "mean": 205.222,
          "n": 3
        },
        "per_stream_decode_tok_s": {
          "min": 123.505,
          "median": 129.307,
          "max": 131.517,
          "mean": 128.709,
          "n": 6
        },
        "per_stream_ttft_ms": {
          "min": 425.2,
          "median": 499.68,
          "max": 569.84,
          "mean": 499.0,
          "n": 6
        },
        "power_mean_w": {
          "min": 115.88,
          "median": 118.57,
          "max": 144.35,
          "mean": 126.27,
          "n": 3
        },
        "energy_j_per_1k_tokens": {
          "min": 563.1,
          "median": 573.75,
          "max": 710.32,
          "mean": 615.72,
          "n": 3
        },
        "temp_max_c": {
          "min": 45.0,
          "median": 45.0,
          "max": 46.0,
          "mean": 45.3,
          "n": 3
        },
        "fan_mean_pct": null,
        "clock_floor_mhz": {
          "min": 2190.0,
          "median": 2190.0,
          "max": 2197.0,
          "mean": 2192.3,
          "n": 3
        }
      },
      "thermal_verdict": "no clock drop during decode (draw held at 96.2% of the 150 W cap); no clock drop during decode; no clock drop during decode"
    },
    {
      "level": 4,
      "contention_gate": {
        "attempts": [
          {
            "window_s": 10.0,
            "samples": 20,
            "cards": {
              "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": {
                "utilization_pct": {
                  "min": 0.0,
                  "median": 0.0,
                  "max": 89.0,
                  "mean": 4.5,
                  "n": 20
                },
                "power_w": {
                  "min": 22.63,
                  "median": 22.7,
                  "max": 144.01,
                  "mean": 32.89,
                  "n": 20
                },
                "card_seen": true
              }
            },
            "max_mean_util_pct": 4.5,
            "all_cards_seen": true,
            "attempt": 1,
            "verdict": "quiet: busiest card mean 4.5% <= 5.0%"
          }
        ],
        "passed": true,
        "cards": [
          "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff"
        ]
      },
      "trials": [
        {
          "level": 4,
          "wall_s": 3.5566,
          "streams": [
            {
              "ttft_ms": 637.77,
              "first_token_was_thinking": false,
              "wall_s": 3.5558,
              "eval_count": 256,
              "eval_duration_ns": 2915897000,
              "decode_tok_s": 87.795,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 152126000,
              "prefill_tok_s": 3523.395,
              "load_duration_ns": 340689561,
              "total_duration_ns": 3552845966,
              "streamed_chunks": 256,
              "response_chars": 1043,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 87.389,
              "decode_tok_s_wall_gross": 71.996,
              "num_ctx_option": 4096,
              "stream": 1
            },
            {
              "ttft_ms": 635.72,
              "first_token_was_thinking": false,
              "wall_s": 3.5557,
              "eval_count": 256,
              "eval_duration_ns": 2917117000,
              "decode_tok_s": 87.758,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 203588000,
              "prefill_tok_s": 2632.768,
              "load_duration_ns": 341024553,
              "total_duration_ns": 3552493368,
              "streamed_chunks": 256,
              "response_chars": 1051,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 87.331,
              "decode_tok_s_wall_gross": 71.998,
              "num_ctx_option": 4096,
              "stream": 2
            },
            {
              "ttft_ms": 634.26,
              "first_token_was_thinking": false,
              "wall_s": 3.5552,
              "eval_count": 256,
              "eval_duration_ns": 2917899000,
              "decode_tok_s": 87.734,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 232770000,
              "prefill_tok_s": 2302.702,
              "load_duration_ns": 367825691,
              "total_duration_ns": 3551905822,
              "streamed_chunks": 256,
              "response_chars": 1051,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 87.302,
              "decode_tok_s_wall_gross": 72.008,
              "num_ctx_option": 4096,
              "stream": 3
            },
            {
              "ttft_ms": 532.97,
              "first_token_was_thinking": false,
              "wall_s": 3.5396,
              "eval_count": 256,
              "eval_duration_ns": 3004423000,
              "decode_tok_s": 85.208,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 192521000,
              "prefill_tok_s": 2784.112,
              "load_duration_ns": 334023546,
              "total_duration_ns": 3536346127,
              "streamed_chunks": 256,
              "response_chars": 1032,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 84.814,
              "decode_tok_s_wall_gross": 72.326,
              "num_ctx_option": 4096,
              "stream": 4
            }
          ],
          "streams_ok": 4,
          "streams_failed": 0,
          "stream_errors": [],
          "trial_void": false,
          "per_stream_decode_tok_s": [
            87.795,
            87.758,
            87.734,
            85.208
          ],
          "per_stream_ttft_ms": [
            637.77,
            635.72,
            634.26,
            532.97
          ],
          "aggregate_sum_tok_s": 348.495,
          "aggregate_wall_tok_s": 287.911,
          "tokens_decoded_total": 1024,
          "per_stream_decode": {
            "min": 85.208,
            "median": 87.746,
            "max": 87.795,
            "mean": 87.124,
            "n": 4
          },
          "per_stream_ttft": {
            "min": 532.97,
            "median": 634.99,
            "max": 637.77,
            "mean": 610.18,
            "n": 4
          },
          "power": {
            "idle_baseline": {
              "seconds": 3.0,
              "samples": 6,
              "power_w": {
                "min": 22.65,
                "median": 22.66,
                "max": 22.68,
                "mean": 22.66,
                "n": 6
              },
              "utilization_pct": {
                "min": 0.0,
                "median": 0.0,
                "max": 0.0,
                "mean": 0.0,
                "n": 6
              },
              "memory_used_mib": {
                "min": 19176.0,
                "median": 19176.0,
                "max": 19176.0,
                "mean": 19176.0,
                "n": 6
              },
              "power_w_per_card": {
                "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": 22.66
              },
              "memory_used_mib_per_card": {
                "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": 19176.0
              },
              "power_w_all_cards": 22.66,
              "ups_whole_box": {
                "load_pct": null,
                "realpower_w": null,
                "realpower_nominal_w": null,
                "status": null,
                "note": "the UPS's reading of EVERYTHING it feeds -- the whole box, not the cards. The card figures beside it are board power from nvidia-smi. Neither is the other."
              }
            },
            "mean_w": 115.34,
            "max_w": 138.83,
            "over_idle_w": 92.68,
            "samples": 13,
            "per_gpu": {
              "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": {
                "index": "0",
                "name": "NVIDIA GeForce RTX 5090 Laptop GPU",
                "power_w": {
                  "min": 22.66,
                  "median": 129.61,
                  "max": 138.83,
                  "mean": 94.01,
                  "n": 13
                },
                "power_limit_w": 150.0,
                "power_limit_field": "enforced.power.limit",
                "power_limit_enforced_w": 150.0,
                "cap": "dynamic-boost",
                "power_posture": "dynamic-boost",
                "power_limit_enforced_stats_w": {
                  "min": 150.0,
                  "median": 150.0,
                  "max": 150.0,
                  "mean": 150.0,
                  "n": 13
                },
                "power_posture_note": "dynamic-boost -- nvidia-powerd running and no cap set: the board floats between its power.default_limit and its power.max_limit against the CPU's draw, DURING the scored run. `power_limit_enforced_stats_w` is that limit reduced over this run's own 2 Hz trace and is the only honest cap cell; `power_limit_enforced_w` is the FIRST sample and is kept only so the two can be compared. A single wattage in a cap column for this posture is a fabrication.",
                "cap_cell": "150 W, flat",
                "cap_cell_note": "the enforced limit read 150 W on every one of 13 samples of this run's 2 Hz trace. Under dynamic-boost that is a FINDING, not a range.",
                "memory_used_mib": {
                  "min": 19176.0,
                  "median": 19176.0,
                  "max": 19176.0,
                  "mean": 19176.0,
                  "n": 13
                },
                "utilization_pct": {
                  "min": 0.0,
                  "median": 86.0,
                  "max": 87.0,
                  "mean": 65.2,
                  "n": 13
                },
                "temperature_c": {
                  "min": 37.0,
                  "median": 45.0,
                  "max": 46.0,
                  "mean": 43.2,
                  "n": 13
                },
                "memory_temperature_c": null,
                "fan_pct": null,
                "pcie_width_current": {
                  "min": 8.0,
                  "median": 8.0,
                  "max": 8.0,
                  "mean": 8.0,
                  "n": 13
                },
                "clocks_sm_mhz": {
                  "min": 1590.0,
                  "median": 1995.0,
                  "max": 2002.0,
                  "mean": 1900.8,
                  "n": 13
                },
                "clocks_mem_mhz": {
                  "min": 14001.0,
                  "median": 14001.0,
                  "max": 14001.0,
                  "mean": 14001.0,
                  "n": 13
                },
                "throttle_sw_power_cap_fraction": 0.7692,
                "throttle_hw_slowdown_fraction": 0.0,
                "throttle_sw_thermal_fraction": 0.0,
                "throttle_hw_thermal_fraction": 0.0,
                "n": 13,
                "busy": {
                  "index": "0",
                  "name": "NVIDIA GeForce RTX 5090 Laptop GPU",
                  "power_w": {
                    "min": 36.56,
                    "median": 138.23,
                    "max": 138.83,
                    "mean": 115.34,
                    "n": 10
                  },
                  "power_limit_w": 150.0,
                  "power_limit_field": "enforced.power.limit",
                  "power_limit_enforced_w": 150.0,
                  "cap": "dynamic-boost",
                  "power_posture": "dynamic-boost",
                  "power_limit_enforced_stats_w": {
                    "min": 150.0,
                    "median": 150.0,
                    "max": 150.0,
                    "mean": 150.0,
                    "n": 10
                  },
                  "power_posture_note": "dynamic-boost -- nvidia-powerd running and no cap set: the board floats between its power.default_limit and its power.max_limit against the CPU's draw, DURING the scored run. `power_limit_enforced_stats_w` is that limit reduced over this run's own 2 Hz trace and is the only honest cap cell; `power_limit_enforced_w` is the FIRST sample and is kept only so the two can be compared. A single wattage in a cap column for this posture is a fabrication.",
                  "cap_cell": "150 W, flat",
                  "cap_cell_note": "the enforced limit read 150 W on every one of 10 samples of this run's 2 Hz trace. Under dynamic-boost that is a FINDING, not a range.",
                  "memory_used_mib": {
                    "min": 19176.0,
                    "median": 19176.0,
                    "max": 19176.0,
                    "mean": 19176.0,
                    "n": 10
                  },
                  "utilization_pct": {
                    "min": 79.0,
                    "median": 86.0,
                    "max": 87.0,
                    "mean": 84.7,
                    "n": 10
                  },
                  "temperature_c": {
                    "min": 43.0,
                    "median": 45.0,
                    "max": 46.0,
                    "mean": 45.0,
                    "n": 10
                  },
                  "memory_temperature_c": null,
                  "fan_pct": null,
                  "pcie_width_current": {
                    "min": 8.0,
                    "median": 8.0,
                    "max": 8.0,
                    "mean": 8.0,
                    "n": 10
                  },
                  "clocks_sm_mhz": {
                    "min": 1942.0,
                    "median": 1995.0,
                    "max": 2002.0,
                    "mean": 1991.0,
                    "n": 10
                  },
                  "clocks_mem_mhz": {
                    "min": 14001.0,
                    "median": 14001.0,
                    "max": 14001.0,
                    "mean": 14001.0,
                    "n": 10
                  },
                  "throttle_sw_power_cap_fraction": 1.0,
                  "throttle_hw_slowdown_fraction": 0.0,
                  "throttle_sw_thermal_fraction": 0.0,
                  "throttle_hw_thermal_fraction": 0.0,
                  "n": 10
                },
                "busy_samples": 10
              }
            }
          },
          "energy_j_per_1k_tokens": 400.61,
          "thermal": {
            "rule_version": 2,
            "measured_over": "samples where the GPU was busy (utilisation > 0)",
            "busy_samples": 10,
            "temp_max_c": 46.0,
            "temp_min_c": 43.0,
            "temp_rise_c": 3.0,
            "memory_temp_max_c": null,
            "memory_temp_median_c": null,
            "memory_temp_readable": false,
            "pcie_width_min": 8.0,
            "pcie_width_max": 8.0,
            "pcie_width_median": 8.0,
            "fan_mean_pct": null,
            "fan_max_pct": null,
            "clock_floor_mhz": 1942.0,
            "clock_max_mhz": 2002.0,
            "clock_median_mhz": 1995.0,
            "mem_clock_floor_mhz": 14001.0,
            "mem_clock_median_mhz": 14001.0,
            "mem_clock_max_mhz": 14001.0,
            "mem_clock_held": true,
            "memory_bus_width_bits": 256,
            "memory_technology": "GDDR7",
            "memory_bits_per_clock": null,
            "peak_mem_bandwidth_gbs_at_median_clock": null,
            "peak_mem_bandwidth_gbs_at_floor_clock": null,
            "peak_mem_bandwidth_gbs_at_max_clock": null,
            "peak_mem_bandwidth_withheld_because": "this harness's peak-bandwidth formula assumes two bits per memory clock per pin, which is exact for GDDR6 and GDDR6X and is checked against three published bandwidths. It has no receipt for the bits-per-clock of this board's memory, so the DERIVED bytes-per-second ceiling is withheld rather than computed from an unverified constant. The memory clock itself is a reading and is published unchanged.",
            "peak_mem_bandwidth_formula": "GB/s = clocks.mem MHz x 1e6 x <bits-per-clock for THIS board's memory> x (bus bits / 8) / 1e9; a CEILING at the clock that was observed, never an achieved rate. The bits-per-clock is registered per board and is NOT defaulted.",
            "clock_dropped": false,
            "power_limit_w": 150.0,
            "power_mean_w": 115.34,
            "power_pct_of_limit": 76.9,
            "power_pinned_at_limit": false,
            "sw_power_cap_active_fraction": 1.0,
            "hw_slowdown_active_fraction": 0.0,
            "sw_thermal_active_fraction": 0.0,
            "hw_thermal_active_fraction": 0.0,
            "at_throttling_temperature": false,
            "stop_core_temp_c": 83.0,
            "stop_memory_temp_c": 100.0,
            "hit_core_temp_stop": false,
            "hit_memory_temp_stop": false,
            "throttle_temperature_c": 80.0,
            "temp_rising_within_run": true,
            "verdict": "no clock drop during decode"
          }
        },
        {
          "level": 4,
          "wall_s": 3.6561,
          "streams": [
            {
              "ttft_ms": 767.36,
              "first_token_was_thinking": false,
              "wall_s": 3.6549,
              "eval_count": 256,
              "eval_duration_ns": 2885276000,
              "decode_tok_s": 88.726,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 154119000,
              "prefill_tok_s": 3477.832,
              "load_duration_ns": 352468743,
              "total_duration_ns": 3652075677,
              "streamed_chunks": 256,
              "response_chars": 1041,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 88.311,
              "decode_tok_s_wall_gross": 70.043,
              "num_ctx_option": 4096,
              "stream": 1
            },
            {
              "ttft_ms": 763.93,
              "first_token_was_thinking": false,
              "wall_s": 3.6553,
              "eval_count": 256,
              "eval_duration_ns": 2888126000,
              "decode_tok_s": 88.639,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 241579000,
              "prefill_tok_s": 2218.736,
              "load_duration_ns": 296859975,
              "total_duration_ns": 3651520578,
              "streamed_chunks": 256,
              "response_chars": 1043,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 88.194,
              "decode_tok_s_wall_gross": 70.036,
              "num_ctx_option": 4096,
              "stream": 2
            },
            {
              "ttft_ms": 766.43,
              "first_token_was_thinking": false,
              "wall_s": 3.6554,
              "eval_count": 256,
              "eval_duration_ns": 2886247000,
              "decode_tok_s": 88.696,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 182797000,
              "prefill_tok_s": 2932.214,
              "load_duration_ns": 308363663,
              "total_duration_ns": 3652444370,
              "streamed_chunks": 256,
              "response_chars": 1044,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 88.268,
              "decode_tok_s_wall_gross": 70.034,
              "num_ctx_option": 4096,
              "stream": 3
            },
            {
              "ttft_ms": 765.7,
              "first_token_was_thinking": false,
              "wall_s": 3.6551,
              "eval_count": 256,
              "eval_duration_ns": 2887136000,
              "decode_tok_s": 88.669,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 212642000,
              "prefill_tok_s": 2520.669,
              "load_duration_ns": 307148381,
              "total_duration_ns": 3651906900,
              "streamed_chunks": 256,
              "response_chars": 1041,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 88.253,
              "decode_tok_s_wall_gross": 70.039,
              "num_ctx_option": 4096,
              "stream": 4
            }
          ],
          "streams_ok": 4,
          "streams_failed": 0,
          "stream_errors": [],
          "trial_void": false,
          "per_stream_decode_tok_s": [
            88.726,
            88.639,
            88.696,
            88.669
          ],
          "per_stream_ttft_ms": [
            767.36,
            763.93,
            766.43,
            765.7
          ],
          "aggregate_sum_tok_s": 354.73,
          "aggregate_wall_tok_s": 280.081,
          "tokens_decoded_total": 1024,
          "per_stream_decode": {
            "min": 88.639,
            "median": 88.683,
            "max": 88.726,
            "mean": 88.683,
            "n": 4
          },
          "per_stream_ttft": {
            "min": 763.93,
            "median": 766.07,
            "max": 767.36,
            "mean": 765.86,
            "n": 4
          },
          "power": {
            "idle_baseline": {
              "seconds": 3.0,
              "samples": 6,
              "power_w": {
                "min": 22.87,
                "median": 30.17,
                "max": 138.04,
                "mean": 53.37,
                "n": 6
              },
              "utilization_pct": {
                "min": 0.0,
                "median": 0.0,
                "max": 87.0,
                "mean": 14.5,
                "n": 6
              },
              "memory_used_mib": {
                "min": 19176.0,
                "median": 19176.0,
                "max": 19176.0,
                "mean": 19176.0,
                "n": 6
              },
              "power_w_per_card": {
                "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": 53.37
              },
              "memory_used_mib_per_card": {
                "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": 19176.0
              },
              "power_w_all_cards": 53.37,
              "ups_whole_box": {
                "load_pct": null,
                "realpower_w": null,
                "realpower_nominal_w": null,
                "status": null,
                "note": "the UPS's reading of EVERYTHING it feeds -- the whole box, not the cards. The card figures beside it are board power from nvidia-smi. Neither is the other."
              }
            },
            "mean_w": 129.38,
            "max_w": 139.32,
            "over_idle_w": 76.01,
            "samples": 14,
            "per_gpu": {
              "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": {
                "index": "0",
                "name": "NVIDIA GeForce RTX 5090 Laptop GPU",
                "power_w": {
                  "min": 22.83,
                  "median": 128.28,
                  "max": 139.32,
                  "mean": 94.07,
                  "n": 14
                },
                "power_limit_w": 150.0,
                "power_limit_field": "enforced.power.limit",
                "power_limit_enforced_w": 150.0,
                "cap": "dynamic-boost",
                "power_posture": "dynamic-boost",
                "power_limit_enforced_stats_w": {
                  "min": 150.0,
                  "median": 150.0,
                  "max": 150.0,
                  "mean": 150.0,
                  "n": 14
                },
                "power_posture_note": "dynamic-boost -- nvidia-powerd running and no cap set: the board floats between its power.default_limit and its power.max_limit against the CPU's draw, DURING the scored run. `power_limit_enforced_stats_w` is that limit reduced over this run's own 2 Hz trace and is the only honest cap cell; `power_limit_enforced_w` is the FIRST sample and is kept only so the two can be compared. A single wattage in a cap column for this posture is a fabrication.",
                "cap_cell": "150 W, flat",
                "cap_cell_note": "the enforced limit read 150 W on every one of 14 samples of this run's 2 Hz trace. Under dynamic-boost that is a FINDING, not a range.",
                "memory_used_mib": {
                  "min": 19176.0,
                  "median": 19178.0,
                  "max": 19178.0,
                  "mean": 19177.6,
                  "n": 14
                },
                "utilization_pct": {
                  "min": 0.0,
                  "median": 86.5,
                  "max": 87.0,
                  "mean": 55.8,
                  "n": 14
                },
                "temperature_c": {
                  "min": 38.0,
                  "median": 45.5,
                  "max": 47.0,
                  "mean": 43.8,
                  "n": 14
                },
                "memory_temperature_c": null,
                "fan_pct": null,
                "pcie_width_current": {
                  "min": 8.0,
                  "median": 8.0,
                  "max": 8.0,
                  "mean": 8.0,
                  "n": 14
                },
                "clocks_sm_mhz": {
                  "min": 1590.0,
                  "median": 1987.0,
                  "max": 2032.0,
                  "mean": 1884.4,
                  "n": 14
                },
                "clocks_mem_mhz": {
                  "min": 14001.0,
                  "median": 14001.0,
                  "max": 14001.0,
                  "mean": 14001.0,
                  "n": 14
                },
                "throttle_sw_power_cap_fraction": 0.6429,
                "throttle_hw_slowdown_fraction": 0.0,
                "throttle_sw_thermal_fraction": 0.0,
                "throttle_hw_thermal_fraction": 0.0,
                "n": 14,
                "busy": {
                  "index": "0",
                  "name": "NVIDIA GeForce RTX 5090 Laptop GPU",
                  "power_w": {
                    "min": 72.71,
                    "median": 139.16,
                    "max": 139.32,
                    "mean": 129.38,
                    "n": 9
                  },
                  "power_limit_w": 150.0,
                  "power_limit_field": "enforced.power.limit",
                  "power_limit_enforced_w": 150.0,
                  "cap": "dynamic-boost",
                  "power_posture": "dynamic-boost",
                  "power_limit_enforced_stats_w": {
                    "min": 150.0,
                    "median": 150.0,
                    "max": 150.0,
                    "mean": 150.0,
                    "n": 9
                  },
                  "power_posture_note": "dynamic-boost -- nvidia-powerd running and no cap set: the board floats between its power.default_limit and its power.max_limit against the CPU's draw, DURING the scored run. `power_limit_enforced_stats_w` is that limit reduced over this run's own 2 Hz trace and is the only honest cap cell; `power_limit_enforced_w` is the FIRST sample and is kept only so the two can be compared. A single wattage in a cap column for this posture is a fabrication.",
                  "cap_cell": "150 W, flat",
                  "cap_cell_note": "the enforced limit read 150 W on every one of 9 samples of this run's 2 Hz trace. Under dynamic-boost that is a FINDING, not a range.",
                  "memory_used_mib": {
                    "min": 19178.0,
                    "median": 19178.0,
                    "max": 19178.0,
                    "mean": 19178.0,
                    "n": 9
                  },
                  "utilization_pct": {
                    "min": 86.0,
                    "median": 87.0,
                    "max": 87.0,
                    "mean": 86.8,
                    "n": 9
                  },
                  "temperature_c": {
                    "min": 45.0,
                    "median": 46.0,
                    "max": 47.0,
                    "mean": 45.9,
                    "n": 9
                  },
                  "memory_temperature_c": null,
                  "fan_pct": null,
                  "pcie_width_current": {
                    "min": 8.0,
                    "median": 8.0,
                    "max": 8.0,
                    "mean": 8.0,
                    "n": 9
                  },
                  "clocks_sm_mhz": {
                    "min": 1987.0,
                    "median": 1995.0,
                    "max": 2032.0,
                    "mean": 1998.0,
                    "n": 9
                  },
                  "clocks_mem_mhz": {
                    "min": 14001.0,
                    "median": 14001.0,
                    "max": 14001.0,
                    "mean": 14001.0,
                    "n": 9
                  },
                  "throttle_sw_power_cap_fraction": 1.0,
                  "throttle_hw_slowdown_fraction": 0.0,
                  "throttle_sw_thermal_fraction": 0.0,
                  "throttle_hw_thermal_fraction": 0.0,
                  "n": 9
                },
                "busy_samples": 9
              }
            }
          },
          "energy_j_per_1k_tokens": 461.94,
          "thermal": {
            "rule_version": 2,
            "measured_over": "samples where the GPU was busy (utilisation > 0)",
            "busy_samples": 9,
            "temp_max_c": 47.0,
            "temp_min_c": 45.0,
            "temp_rise_c": 2.0,
            "memory_temp_max_c": null,
            "memory_temp_median_c": null,
            "memory_temp_readable": false,
            "pcie_width_min": 8.0,
            "pcie_width_max": 8.0,
            "pcie_width_median": 8.0,
            "fan_mean_pct": null,
            "fan_max_pct": null,
            "clock_floor_mhz": 1987.0,
            "clock_max_mhz": 2032.0,
            "clock_median_mhz": 1995.0,
            "mem_clock_floor_mhz": 14001.0,
            "mem_clock_median_mhz": 14001.0,
            "mem_clock_max_mhz": 14001.0,
            "mem_clock_held": true,
            "memory_bus_width_bits": 256,
            "memory_technology": "GDDR7",
            "memory_bits_per_clock": null,
            "peak_mem_bandwidth_gbs_at_median_clock": null,
            "peak_mem_bandwidth_gbs_at_floor_clock": null,
            "peak_mem_bandwidth_gbs_at_max_clock": null,
            "peak_mem_bandwidth_withheld_because": "this harness's peak-bandwidth formula assumes two bits per memory clock per pin, which is exact for GDDR6 and GDDR6X and is checked against three published bandwidths. It has no receipt for the bits-per-clock of this board's memory, so the DERIVED bytes-per-second ceiling is withheld rather than computed from an unverified constant. The memory clock itself is a reading and is published unchanged.",
            "peak_mem_bandwidth_formula": "GB/s = clocks.mem MHz x 1e6 x <bits-per-clock for THIS board's memory> x (bus bits / 8) / 1e9; a CEILING at the clock that was observed, never an achieved rate. The bits-per-clock is registered per board and is NOT defaulted.",
            "clock_dropped": false,
            "power_limit_w": 150.0,
            "power_mean_w": 129.38,
            "power_pct_of_limit": 86.3,
            "power_pinned_at_limit": false,
            "sw_power_cap_active_fraction": 1.0,
            "hw_slowdown_active_fraction": 0.0,
            "sw_thermal_active_fraction": 0.0,
            "hw_thermal_active_fraction": 0.0,
            "at_throttling_temperature": false,
            "stop_core_temp_c": 83.0,
            "stop_memory_temp_c": 100.0,
            "hit_core_temp_stop": false,
            "hit_memory_temp_stop": false,
            "throttle_temperature_c": 80.0,
            "temp_rising_within_run": false,
            "verdict": "no clock drop during decode"
          }
        },
        {
          "level": 4,
          "wall_s": 3.6123,
          "streams": [
            {
              "ttft_ms": 697.23,
              "first_token_was_thinking": false,
              "wall_s": 3.6105,
              "eval_count": 256,
              "eval_duration_ns": 2911709000,
              "decode_tok_s": 87.921,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 172257000,
              "prefill_tok_s": 3111.63,
              "load_duration_ns": 360692646,
              "total_duration_ns": 3608659178,
              "streamed_chunks": 256,
              "response_chars": 1044,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 87.531,
              "decode_tok_s_wall_gross": 70.905,
              "num_ctx_option": 4096,
              "stream": 1
            },
            {
              "ttft_ms": 698.93,
              "first_token_was_thinking": false,
              "wall_s": 3.6115,
              "eval_count": 256,
              "eval_duration_ns": 2910931000,
              "decode_tok_s": 87.944,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 143586000,
              "prefill_tok_s": 3732.954,
              "load_duration_ns": 367841350,
              "total_duration_ns": 3609379851,
              "streamed_chunks": 256,
              "response_chars": 1041,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 87.55,
              "decode_tok_s_wall_gross": 70.884,
              "num_ctx_option": 4096,
              "stream": 2
            },
            {
              "ttft_ms": 696.83,
              "first_token_was_thinking": false,
              "wall_s": 3.6115,
              "eval_count": 256,
              "eval_duration_ns": 2911635000,
              "decode_tok_s": 87.923,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 231577000,
              "prefill_tok_s": 2314.565,
              "load_duration_ns": 316571314,
              "total_duration_ns": 3607975684,
              "streamed_chunks": 256,
              "response_chars": 1043,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 87.489,
              "decode_tok_s_wall_gross": 70.885,
              "num_ctx_option": 4096,
              "stream": 3
            },
            {
              "ttft_ms": 697.61,
              "first_token_was_thinking": false,
              "wall_s": 3.6115,
              "eval_count": 256,
              "eval_duration_ns": 2911668000,
              "decode_tok_s": 87.922,
              "prompt_eval_count": 536,
              "prompt_eval_duration_ns": 202287000,
              "prefill_tok_s": 2649.701,
              "load_duration_ns": 319182185,
              "total_duration_ns": 3608201098,
              "streamed_chunks": 256,
              "response_chars": 1041,
              "thinking_chars": 0,
              "done_reason": "length",
              "decode_tok_s_wall": 87.513,
              "decode_tok_s_wall_gross": 70.886,
              "num_ctx_option": 4096,
              "stream": 4
            }
          ],
          "streams_ok": 4,
          "streams_failed": 0,
          "stream_errors": [],
          "trial_void": false,
          "per_stream_decode_tok_s": [
            87.921,
            87.944,
            87.923,
            87.922
          ],
          "per_stream_ttft_ms": [
            697.23,
            698.93,
            696.83,
            697.61
          ],
          "aggregate_sum_tok_s": 351.71,
          "aggregate_wall_tok_s": 283.478,
          "tokens_decoded_total": 1024,
          "per_stream_decode": {
            "min": 87.921,
            "median": 87.922,
            "max": 87.944,
            "mean": 87.928,
            "n": 4
          },
          "per_stream_ttft": {
            "min": 696.83,
            "median": 697.42,
            "max": 698.93,
            "mean": 697.65,
            "n": 4
          },
          "power": {
            "idle_baseline": {
              "seconds": 3.0,
              "samples": 6,
              "power_w": {
                "min": 22.76,
                "median": 25.55,
                "max": 139.16,
                "mean": 54.05,
                "n": 6
              },
              "utilization_pct": {
                "min": 0.0,
                "median": 0.0,
                "max": 87.0,
                "mean": 29.0,
                "n": 6
              },
              "memory_used_mib": {
                "min": 19178.0,
                "median": 19178.0,
                "max": 19178.0,
                "mean": 19178.0,
                "n": 6
              },
              "power_w_per_card": {
                "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": 54.05
              },
              "memory_used_mib_per_card": {
                "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": 19178.0
              },
              "power_w_all_cards": 54.05,
              "ups_whole_box": {
                "load_pct": null,
                "realpower_w": null,
                "realpower_nominal_w": null,
                "status": null,
                "note": "the UPS's reading of EVERYTHING it feeds -- the whole box, not the cards. The card figures beside it are board power from nvidia-smi. Neither is the other."
              }
            },
            "mean_w": 136.15,
            "max_w": 138.87,
            "over_idle_w": 82.1,
            "samples": 14,
            "per_gpu": {
              "GPU-edff232c-7dbf-2bac-07fb-921a7c9eecff": {
                "index": "0",
                "name": "NVIDIA GeForce RTX 5090 Laptop GPU",
                "power_w": {
                  "min": 22.75,
                  "median": 128.61,
                  "max": 138.87,
                  "mean": 95.41,
                  "n": 14
                },
                "power_limit_w": 150.0,
                "power_limit_field": "enforced.power.limit",
                "power_limit_enforced_w": 150.0,
                "cap": "dynamic-boost",
                "power_posture": "dynamic-boost",
                "power_limit_enforced_stats_w": {
                  "min": 150.0,
                  "median": 150.0,
                  "max": 150.0,
                  "mean": 150.0,
                  "n": 14
                },
                "power_posture_note": "dynamic-boost -- nvidia-powerd running and no cap set: the board floats between its power.default_limit and its power.max_limit against the CPU's draw, DURING the scored run. `power_limit_enforced_stats_w` is that limit reduced over this run's own 2 Hz trace and is the only honest cap cell; `power_limit_enforced_w` is the FIRST sample and is kept only so the two can be compared. A single wattage in a cap column for this posture is a fabrication.",
                "cap_cell": "150 W, flat",
                "cap_cell_note": "the enforced limit read 150 W on every one of 14 samples of this run's 2 Hz trace. Under dynamic-boost that is a FINDING, not a range.",
                "memory_used_mib": {
                  "min": 19178.0,
                  "median": 19178.0,
                  "max": 19178.0,
                  "mean": 19178.0,
                  "n": 14
                },
                "utilization_pct": {
                  "min": 0.0,
                  "median": 85.0,
                  "max": 87.0,
                  "mean": 49.1,
                  "n": 14
                },
                "temperature_c": {
                  "min": 39.0,
                  "median": 46.0,
                  "max": 47.0,
                  "mean": 44.7,
                  "n": 14
                },
                "memory_temperature_c": null,
                "fan_pct": null,
                "pcie_width_current": {
                  "min": 8.0,
                  "median": 8.0,
                  "max": 8.0,
                  "mean": 8.0,
                  "n": 14
                },
                "clocks_sm_mhz": {
                  "min": 1590.0,
                  "median": 1991.0,
                  "max": 2032.0,
                  "mean": 1903.6,
                  "n": 14
                },
                "clocks_mem_mhz": {
                  "min": 14001.0,
                  "median": 14001.0,
                  "max": 14001.0,
                  "mean": 14001.0,
                  "n": 14
                },
                "throttle_sw_power_cap_fraction": 0.7143,
                "throttle_hw_slowdown_fraction": 0.0,
                "throttle_sw_thermal_fraction": 0.0,
                "throttle_hw_thermal_fraction": 0.0,
                "n": 14,
                "busy": {
                  "index": "0",
                  "name": "NVIDIA GeForce RTX 5090 Laptop GPU",
                  "power_w": {
                    "min": 118.81,
                    "median": 138.57,
                    "max": 138.87,
                    "mean": 136.15,
                    "n": 8
                  },
                  "power_limit_w": 150.0,
                  "power_limit_field": "enforced.power.limit",
                  "power_limit_enforced_w": 150.0,
                  "cap": "dynamic-boost",
                  "power_posture": "dynamic-boost",
                  "power_limit_enforced_stats_w": {
                    "min": 150.0,
                    "median": 150.0,
                    "max": 150.0,
                    "mean": 150.0,
                    "n": 8
                  },
                  "power_posture_note": "dynamic-boost -- nvidia-powerd running and no cap set: the board floats between its power.default_limit and its power.max_limit against the CPU's draw, DURING the scored run. `power_limit_enforced_stats_w` is that limit reduced over this run's own 2 Hz trace and is the only honest cap cell; `power_limit_enforced_w` is the FIRST sample and is kept only so the two can be compared. A single wattage in a cap column for this posture is a fabrication.",
                  "cap_cell": "150 W, flat",
                  "cap_cell_note": "the enforced limit read 150 W on every one of 8 samples of this run's 2 Hz trace. Under dynamic-boost that is a FINDING, not a range.",
                  "memory_used_mib": {
                    "min": 19178.0,
                    "median": 19178.0,
                    "max": 19178.0,
                    "mean": 19178.0,
                    "n": 8
                  },
                  "utilization_pct": {
                    "min": 85.0,
                    "median": 86.0,
                    "max": 87.0,
                    "mean": 86.0,
                    "n": 8
                  },
                  "temperature_c": {
                    "min": 46.0,
                    "median": 47.0,
                    "max": 47.0,
                    "mean": 46.8,
                    "n": 8
                  },
                  "memory_temperature_c": null,
                  "fan_pct": null,
                  "pcie_width_current": {
                    "min": 8.0,
                    "median": 8.0,
                    "max": 8.0,
                    "mean": 8.0,
                    "n": 8
                  },
                  "clocks_sm_mhz": {
                    "min": 1987.0,
                    "median": 1995.0,
                    "max": 2002.0,
                    "mean": 1993.8,
                    "n": 8
                  },
                  "clocks_mem_mhz": {
                    "min": 14001.0,
                    "median": 14001.0,
                    "max": 14001.0,
                    "mean": 14001.0,
                    "n": 8
                  },
                  "throttle_sw_power_cap_fraction": 1.0,
                  "throttle_hw_slowdown_fraction": 0.0,
                  "throttle_sw_thermal_fraction": 0.0,
                  "throttle_hw_thermal_fraction": 0.0,
                  "n": 8
                },
                "busy_samples": 8
              }
            }
          },
          "energy_j_per_1k_tokens": 480.28,
          "thermal": {
            "rule_version": 2,
            "measured_over": "samples where the GPU was busy (utilisation > 0)",
            "busy_samples": 8,
            "temp_max_c": 47.0,
            "temp_min_c": 46.0,
            "temp_rise_c": 1.0,
            "memory_temp_max_c": null,
            "memory_temp_median_c": null,
            "memory_temp_readable": false,
            "pcie_width_min": 8.0,
            "pcie_width_max": 8.0,
            "pcie_width_median": 8.0,
            "fan_mean_pct": null,
            "fan_max_pct": null,
            "clock_floor_mhz": 1987.0,
            "clock_max_mhz": 2002.0,
            "clock_median_mhz": 1995.0,
            "mem_clock_floor_mhz": 14001.0,
            "mem_clock_median_mhz": 14001.0,
            "mem_clock_max_mhz": 14001.0,
            "mem_clock_held": true,
            "memory_bus_width_bits": 256,
            "memory_technology": "GDDR7",
            "memory_bits_per_clock": null,
            "peak_mem_bandwidth_gbs_at_median_clock": null,
            "peak_mem_bandwidth_gbs_at_floor_clock": null,
            "peak_mem_bandwidth_gbs_at_max_clock": null,
            "peak_mem_bandwidth_withheld_because": "this harness's peak-bandwidth formula assumes two bits per memory clock per pin, which is exact for GDDR6 and GDDR6X and is checked against three published bandwidths. It has no receipt for the bits-per-clock of this board's memory, so the DERIVED bytes-per-second ceiling is withheld rather than computed from an unverified constant. The memory clock itself is a reading and is published unchanged.",
            "peak_mem_bandwidth_formula": "GB/s = clocks.mem MHz x 1e6 x <bits-per-clock for THIS board's memory> x (bus bits / 8) / 1e9; a CEILING at the clock that was observed, never an achieved rate. The bits-per-clock is registered per board and is NOT defaulted.",
            "clock_dropped": false,
            "power_limit_w": 150.0,
            "power_mean_w": 136.15,
            "power_pct_of_limit": 90.8,
            "power_pinned_at_limit": true,
            "sw_power_cap_active_fraction": 1.0,
            "hw_slowdown_active_fraction": 0.0,
            "sw_thermal_active_fraction": 0.0,
            "hw_thermal_active_fraction": 0.0,
            "at_throttling_temperature": false,
            "stop_core_temp_c": 83.0,
            "stop_memory_temp_c": 100.0,
            "hit_core_temp_stop": false,
            "hit_memory_temp_stop": false,
            "throttle_temperature_c": 80.0,
            "temp_rising_within_run": false,
            "verdict": "no clock drop during decode (draw held at 90.8% of the 150 W cap)"
          }
        }
      ],
      "spread_gate": {
        "limit": 0.15,
        "values": [
          348.495,
          354.73,
          351.71
        ],
        "spread_fraction": 0.0177,
        "verdict": "spread 1.8% of the median within the 15% gate"
      },
      "voided_by_gate": false,
      "derived": {
        "aggregate_sum_tok_s": {
          "min": 348.495,
          "median": 351.71,
          "max": 354.73,
          "mean": 351.645,
          "n": 3
        },
        "aggregate_wall_tok_s": {
          "min": 280.081,
          "median": 283.478,
          "max": 287.911,
          "mean": 283.823,
          "n": 3
        },
        "per_stream_decode_tok_s": {
          "min": 85.208,
          "median": 87.922,
          "max": 88.726,
          "mean": 87.911,
          "n": 12
        },
        "per_stream_ttft_ms": {
          "min": 532.97,
          "median": 697.42,
          "max": 767.36,
          "mean": 691.23,
          "n": 12
        },
        "power_mean_w": {
          "min": 115.34,
          "median": 129.38,
          "max": 136.15,
          "mean": 126.96,
          "n": 3
        },
        "energy_j_per_1k_tokens": {
          "min": 400.61,
          "median": 461.94,
          "max": 480.28,
          "mean": 447.61,
          "n": 3
        },
        "temp_max_c": {
          "min": 46.0,
          "median": 47.0,
          "max": 47.0,
          "mean": 46.7,
          "n": 3
        },
        "fan_mean_pct": null,
        "clock_floor_mhz": {
          "min": 1942.0,
          "median": 1987.0,
          "max": 1987.0,
          "mean": 1972.0,
          "n": 3
        }
      },
      "thermal_verdict": "no clock drop during decode; no clock drop during decode; no clock drop during decode (draw held at 90.8% of the 150 W cap)"
    }
  ],
  "errors": [],
  "aggregate_note": "aggregate_sum_tok_s is the sum of the streams' own decode rates -- what the card produced while producing. aggregate_wall_tok_s is every stream's tokens over the wall time of the whole level -- what a caller experiences, queueing and ragged finish included. Both are printed; neither is the 'real' one on its own.",
  "ollama_version": "0.32.13",
  "think_mode": false,
  "think_note": "think:false accepted",
  "ps_after_load": [
    {
      "name": "gemma4:26b",
      "model": "gemma4:26b",
      "size": 17728054230,
      "digest": "5571076f3d70050487b26b341705799e0ab29b808164f90d20d4cf84f699d251",
      "details": {
        "parent_model": "",
        "format": "gguf",
        "family": "gemma4",
        "families": [
          "gemma4"
        ],
        "parameter_size": "25.8B",
        "quantization_level": "Q4_K_M"
      },
      "expires_at": "2026-09-21T20:59:58.22077324Z",
      "size_vram": 17728054230,
      "context_length": 4096
    }
  ],
  "size_total": 17728054230,
  "size_vram": 17728054230,
  "fit_verdict": "fits: fully resident on the card",
  "ps_final": [],
  "finished_utc": "2026-09-21T20:51:29Z",
  "headline": "one-dynboost takes 4 parallel streams at 351.7 tok/s aggregate (sum of streams) / 283.5 tok/s over wall"
}