{
  "schema_version": 2,
  "title": "Meet Jev — character-design v2",
  "condition_name": "character-design v2",
  "scoring_regime": "Actual returned five-choice categories; probability diagnostics are audit flags, NOT personality probabilities.",
  "instrument": "Johnson (2014) IPIP-NEO-120",
  "target": "Designed tendencies of the existing fictional human character, NOT discovered human experiences or actual AI psychology",
  "source_attribution": "Johnson, J. A. (2014). Journal of Research in Personality, 51, 78–89. doi:10.1016/j.jrp.2014.05.003; International Personality Item Pool, ipip.ori.org",
  "rights": "IPIP items/scales/inventories public domain; not the proprietary NEO-PI-R; not the copyrighted journal article.",
  "source_discrepancies": [
    {
      "number": 58,
      "official_key": "Believe that there is no absolute right and wrong.",
      "paper_table": "Believe that there is no absolute right or wrong",
      "decision": "Use exact official IPIP Johnson key wording (and), not paper-table or. Disclosed source discrepancy; no mixing with Maples."
    }
  ],
  "scale_display": "25 × (mean keyed category − 1); 0–100 fictional design scale position, not human percentile, probability, IQ or diagnosis",
  "coverage_rule": "Primary requires every item of the group in each of three passes; no imputation. Strict sensitivity separately uses original common-item thresholds.",
  "domains": [
    {
      "code": "N",
      "name": "Neuroticism",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 24,
      "common_items": 24,
      "common_item_ids": [
        "ipip001",
        "ipip006",
        "ipip011",
        "ipip016",
        "ipip021",
        "ipip026",
        "ipip031",
        "ipip036",
        "ipip041",
        "ipip046",
        "ipip051",
        "ipip056",
        "ipip061",
        "ipip066",
        "ipip071",
        "ipip076",
        "ipip081",
        "ipip086",
        "ipip091",
        "ipip096",
        "ipip101",
        "ipip106",
        "ipip111",
        "ipip116"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 24,
          "missing_count": 0,
          "observed_sum": 38,
          "available_mean": 1.5833333333333333,
          "available_scale_position": 14.583333333333332,
          "official_complete_raw_total": 38,
          "full_raw_total_bounds": [
            38,
            38
          ],
          "full_scale_position_bounds": [
            14.583333333333332,
            14.583333333333332
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 24,
          "missing_count": 0,
          "observed_sum": 38,
          "available_mean": 1.5833333333333333,
          "available_scale_position": 14.583333333333332,
          "official_complete_raw_total": 38,
          "full_raw_total_bounds": [
            38,
            38
          ],
          "full_scale_position_bounds": [
            14.583333333333332,
            14.583333333333332
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 24,
          "missing_count": 0,
          "observed_sum": 38,
          "available_mean": 1.5833333333333333,
          "available_scale_position": 14.583333333333332,
          "official_complete_raw_total": 38,
          "full_raw_total_bounds": [
            38,
            38
          ],
          "full_scale_position_bounds": [
            14.583333333333332,
            14.583333333333332
          ]
        }
      ],
      "common_mean_by_pass": [
        1.5833333333333333,
        1.5833333333333333,
        1.5833333333333333
      ],
      "common_position_by_pass": [
        14.583333333333332,
        14.583333333333332,
        14.583333333333332
      ],
      "display_mean": 1.5833333333333333,
      "display_scale_position": 14.583333333333332,
      "repeat_statistics": {
        "n": 3,
        "mean": 14.583333333333332,
        "sample_sd": 0.0,
        "min": 14.583333333333332,
        "max": 14.583333333333332,
        "mean_minus_sd": 14.583333333333332
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 24,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": false,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 24,
        "common_items": 20,
        "common_item_ids": [
          "ipip001",
          "ipip006",
          "ipip011",
          "ipip016",
          "ipip026",
          "ipip036",
          "ipip046",
          "ipip051",
          "ipip056",
          "ipip061",
          "ipip066",
          "ipip076",
          "ipip081",
          "ipip086",
          "ipip091",
          "ipip096",
          "ipip101",
          "ipip106",
          "ipip111",
          "ipip116"
        ],
        "omitted_item_ids": [
          "ipip021",
          "ipip031",
          "ipip041",
          "ipip071"
        ],
        "complete_all_passes": false,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 23,
            "missing_count": 1,
            "observed_sum": 37,
            "available_mean": 1.608695652173913,
            "available_scale_position": 15.217391304347828,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              38,
              42
            ],
            "full_scale_position_bounds": [
              14.583333333333332,
              18.75
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 22,
            "missing_count": 2,
            "observed_sum": 36,
            "available_mean": 1.6363636363636365,
            "available_scale_position": 15.909090909090912,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              38,
              46
            ],
            "full_scale_position_bounds": [
              14.583333333333332,
              22.916666666666668
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 23,
            "missing_count": 1,
            "observed_sum": 37,
            "available_mean": 1.608695652173913,
            "available_scale_position": 15.217391304347828,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              38,
              42
            ],
            "full_scale_position_bounds": [
              14.583333333333332,
              18.75
            ]
          }
        ],
        "common_mean_by_pass": [
          1.7,
          1.7,
          1.7
        ],
        "common_position_by_pass": [
          17.5,
          17.5,
          17.5
        ],
        "display_mean": null,
        "display_scale_position": null,
        "repeat_statistics": {
          "n": 0,
          "mean": null,
          "sample_sd": null,
          "min": null,
          "max": null,
          "mean_minus_sd": null
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 20,
        "constituent_facets_covered": false
      },
      "categorical_minus_strict_position": null
    },
    {
      "code": "E",
      "name": "Extraversion",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 24,
      "common_items": 24,
      "common_item_ids": [
        "ipip002",
        "ipip007",
        "ipip012",
        "ipip017",
        "ipip022",
        "ipip027",
        "ipip032",
        "ipip037",
        "ipip042",
        "ipip047",
        "ipip052",
        "ipip057",
        "ipip062",
        "ipip067",
        "ipip072",
        "ipip077",
        "ipip082",
        "ipip087",
        "ipip092",
        "ipip097",
        "ipip102",
        "ipip107",
        "ipip112",
        "ipip117"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 24,
          "missing_count": 0,
          "observed_sum": 84,
          "available_mean": 3.5,
          "available_scale_position": 62.5,
          "official_complete_raw_total": 84,
          "full_raw_total_bounds": [
            84,
            84
          ],
          "full_scale_position_bounds": [
            62.5,
            62.5
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 24,
          "missing_count": 0,
          "observed_sum": 86,
          "available_mean": 3.5833333333333335,
          "available_scale_position": 64.58333333333334,
          "official_complete_raw_total": 86,
          "full_raw_total_bounds": [
            86,
            86
          ],
          "full_scale_position_bounds": [
            64.58333333333334,
            64.58333333333334
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 24,
          "missing_count": 0,
          "observed_sum": 84,
          "available_mean": 3.5,
          "available_scale_position": 62.5,
          "official_complete_raw_total": 84,
          "full_raw_total_bounds": [
            84,
            84
          ],
          "full_scale_position_bounds": [
            62.5,
            62.5
          ]
        }
      ],
      "common_mean_by_pass": [
        3.5,
        3.5833333333333335,
        3.5
      ],
      "common_position_by_pass": [
        62.5,
        64.58333333333334,
        62.5
      ],
      "display_mean": 3.5277777777777777,
      "display_scale_position": 63.19444444444445,
      "repeat_statistics": {
        "n": 3,
        "mean": 63.19444444444445,
        "sample_sd": 1.2028130608117258,
        "min": 62.5,
        "max": 64.58333333333334,
        "mean_minus_sd": 61.99163138363272
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 24,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 24,
        "common_items": 21,
        "common_item_ids": [
          "ipip002",
          "ipip007",
          "ipip012",
          "ipip017",
          "ipip022",
          "ipip027",
          "ipip032",
          "ipip042",
          "ipip052",
          "ipip057",
          "ipip062",
          "ipip067",
          "ipip072",
          "ipip077",
          "ipip082",
          "ipip087",
          "ipip092",
          "ipip097",
          "ipip102",
          "ipip107",
          "ipip112"
        ],
        "omitted_item_ids": [
          "ipip037",
          "ipip047",
          "ipip117"
        ],
        "complete_all_passes": false,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 22,
            "missing_count": 2,
            "observed_sum": 76,
            "available_mean": 3.4545454545454546,
            "available_scale_position": 61.36363636363637,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              78,
              86
            ],
            "full_scale_position_bounds": [
              56.25,
              64.58333333333334
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 22,
            "missing_count": 2,
            "observed_sum": 80,
            "available_mean": 3.6363636363636362,
            "available_scale_position": 65.9090909090909,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              82,
              90
            ],
            "full_scale_position_bounds": [
              60.416666666666664,
              68.75
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 24,
            "missing_count": 0,
            "observed_sum": 84,
            "available_mean": 3.5,
            "available_scale_position": 62.5,
            "official_complete_raw_total": 84,
            "full_raw_total_bounds": [
              84,
              84
            ],
            "full_scale_position_bounds": [
              62.5,
              62.5
            ]
          }
        ],
        "common_mean_by_pass": [
          3.5238095238095237,
          3.619047619047619,
          3.5238095238095237
        ],
        "common_position_by_pass": [
          63.095238095238095,
          65.47619047619048,
          63.095238095238095
        ],
        "display_mean": 3.5555555555555554,
        "display_scale_position": 63.88888888888889,
        "repeat_statistics": {
          "n": 3,
          "mean": 63.88888888888889,
          "sample_sd": 1.3746434980705409,
          "min": 63.095238095238095,
          "max": 65.47619047619048,
          "mean_minus_sd": 62.51424539081835
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 20,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": -0.6944444444444429
    },
    {
      "code": "O",
      "name": "Openness to Experience",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 24,
      "common_items": 24,
      "common_item_ids": [
        "ipip003",
        "ipip008",
        "ipip013",
        "ipip018",
        "ipip023",
        "ipip028",
        "ipip033",
        "ipip038",
        "ipip043",
        "ipip048",
        "ipip053",
        "ipip058",
        "ipip063",
        "ipip068",
        "ipip073",
        "ipip078",
        "ipip083",
        "ipip088",
        "ipip093",
        "ipip098",
        "ipip103",
        "ipip108",
        "ipip113",
        "ipip118"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 24,
          "missing_count": 0,
          "observed_sum": 103,
          "available_mean": 4.291666666666667,
          "available_scale_position": 82.29166666666667,
          "official_complete_raw_total": 103,
          "full_raw_total_bounds": [
            103,
            103
          ],
          "full_scale_position_bounds": [
            82.29166666666667,
            82.29166666666667
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 24,
          "missing_count": 0,
          "observed_sum": 102,
          "available_mean": 4.25,
          "available_scale_position": 81.25,
          "official_complete_raw_total": 102,
          "full_raw_total_bounds": [
            102,
            102
          ],
          "full_scale_position_bounds": [
            81.25,
            81.25
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 24,
          "missing_count": 0,
          "observed_sum": 101,
          "available_mean": 4.208333333333333,
          "available_scale_position": 80.20833333333333,
          "official_complete_raw_total": 101,
          "full_raw_total_bounds": [
            101,
            101
          ],
          "full_scale_position_bounds": [
            80.20833333333333,
            80.20833333333333
          ]
        }
      ],
      "common_mean_by_pass": [
        4.291666666666667,
        4.25,
        4.208333333333333
      ],
      "common_position_by_pass": [
        82.29166666666667,
        81.25,
        80.20833333333333
      ],
      "display_mean": 4.25,
      "display_scale_position": 81.25,
      "repeat_statistics": {
        "n": 3,
        "mean": 81.25,
        "sample_sd": 1.0416666666666714,
        "min": 80.20833333333333,
        "max": 82.29166666666667,
        "mean_minus_sd": 80.20833333333333
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 24,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 24,
        "common_items": 23,
        "common_item_ids": [
          "ipip003",
          "ipip008",
          "ipip013",
          "ipip018",
          "ipip023",
          "ipip028",
          "ipip033",
          "ipip038",
          "ipip043",
          "ipip048",
          "ipip053",
          "ipip058",
          "ipip063",
          "ipip068",
          "ipip073",
          "ipip078",
          "ipip083",
          "ipip088",
          "ipip093",
          "ipip098",
          "ipip103",
          "ipip113",
          "ipip118"
        ],
        "omitted_item_ids": [
          "ipip108"
        ],
        "complete_all_passes": false,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 24,
            "missing_count": 0,
            "observed_sum": 103,
            "available_mean": 4.291666666666667,
            "available_scale_position": 82.29166666666667,
            "official_complete_raw_total": 103,
            "full_raw_total_bounds": [
              103,
              103
            ],
            "full_scale_position_bounds": [
              82.29166666666667,
              82.29166666666667
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 23,
            "missing_count": 1,
            "observed_sum": 98,
            "available_mean": 4.260869565217392,
            "available_scale_position": 81.5217391304348,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              99,
              103
            ],
            "full_scale_position_bounds": [
              78.125,
              82.29166666666667
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 24,
            "missing_count": 0,
            "observed_sum": 101,
            "available_mean": 4.208333333333333,
            "available_scale_position": 80.20833333333333,
            "official_complete_raw_total": 101,
            "full_raw_total_bounds": [
              101,
              101
            ],
            "full_scale_position_bounds": [
              80.20833333333333,
              80.20833333333333
            ]
          }
        ],
        "common_mean_by_pass": [
          4.260869565217392,
          4.260869565217392,
          4.217391304347826
        ],
        "common_position_by_pass": [
          81.5217391304348,
          81.5217391304348,
          80.43478260869566
        ],
        "display_mean": 4.246376811594203,
        "display_scale_position": 81.15942028985508,
        "repeat_statistics": {
          "n": 3,
          "mean": 81.15942028985508,
          "sample_sd": 0.6275546404235116,
          "min": 80.43478260869566,
          "max": 81.5217391304348,
          "mean_minus_sd": 80.53186564943157
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 20,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0.09057971014492239
    },
    {
      "code": "A",
      "name": "Agreeableness",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 24,
      "common_items": 24,
      "common_item_ids": [
        "ipip004",
        "ipip009",
        "ipip014",
        "ipip019",
        "ipip024",
        "ipip029",
        "ipip034",
        "ipip039",
        "ipip044",
        "ipip049",
        "ipip054",
        "ipip059",
        "ipip064",
        "ipip069",
        "ipip074",
        "ipip079",
        "ipip084",
        "ipip089",
        "ipip094",
        "ipip099",
        "ipip104",
        "ipip109",
        "ipip114",
        "ipip119"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 24,
          "missing_count": 0,
          "observed_sum": 109,
          "available_mean": 4.541666666666667,
          "available_scale_position": 88.54166666666667,
          "official_complete_raw_total": 109,
          "full_raw_total_bounds": [
            109,
            109
          ],
          "full_scale_position_bounds": [
            88.54166666666667,
            88.54166666666667
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 24,
          "missing_count": 0,
          "observed_sum": 109,
          "available_mean": 4.541666666666667,
          "available_scale_position": 88.54166666666667,
          "official_complete_raw_total": 109,
          "full_raw_total_bounds": [
            109,
            109
          ],
          "full_scale_position_bounds": [
            88.54166666666667,
            88.54166666666667
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 24,
          "missing_count": 0,
          "observed_sum": 109,
          "available_mean": 4.541666666666667,
          "available_scale_position": 88.54166666666667,
          "official_complete_raw_total": 109,
          "full_raw_total_bounds": [
            109,
            109
          ],
          "full_scale_position_bounds": [
            88.54166666666667,
            88.54166666666667
          ]
        }
      ],
      "common_mean_by_pass": [
        4.541666666666667,
        4.541666666666667,
        4.541666666666667
      ],
      "common_position_by_pass": [
        88.54166666666667,
        88.54166666666667,
        88.54166666666667
      ],
      "display_mean": 4.541666666666667,
      "display_scale_position": 88.54166666666667,
      "repeat_statistics": {
        "n": 3,
        "mean": 88.54166666666667,
        "sample_sd": 0.0,
        "min": 88.54166666666667,
        "max": 88.54166666666667,
        "mean_minus_sd": 88.54166666666667
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 24,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 24,
        "common_items": 23,
        "common_item_ids": [
          "ipip009",
          "ipip014",
          "ipip019",
          "ipip024",
          "ipip029",
          "ipip034",
          "ipip039",
          "ipip044",
          "ipip049",
          "ipip054",
          "ipip059",
          "ipip064",
          "ipip069",
          "ipip074",
          "ipip079",
          "ipip084",
          "ipip089",
          "ipip094",
          "ipip099",
          "ipip104",
          "ipip109",
          "ipip114",
          "ipip119"
        ],
        "omitted_item_ids": [
          "ipip004"
        ],
        "complete_all_passes": false,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 24,
            "missing_count": 0,
            "observed_sum": 109,
            "available_mean": 4.541666666666667,
            "available_scale_position": 88.54166666666667,
            "official_complete_raw_total": 109,
            "full_raw_total_bounds": [
              109,
              109
            ],
            "full_scale_position_bounds": [
              88.54166666666667,
              88.54166666666667
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 24,
            "missing_count": 0,
            "observed_sum": 109,
            "available_mean": 4.541666666666667,
            "available_scale_position": 88.54166666666667,
            "official_complete_raw_total": 109,
            "full_raw_total_bounds": [
              109,
              109
            ],
            "full_scale_position_bounds": [
              88.54166666666667,
              88.54166666666667
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 23,
            "missing_count": 1,
            "observed_sum": 105,
            "available_mean": 4.565217391304348,
            "available_scale_position": 89.13043478260869,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              106,
              110
            ],
            "full_scale_position_bounds": [
              85.41666666666667,
              89.58333333333333
            ]
          }
        ],
        "common_mean_by_pass": [
          4.565217391304348,
          4.565217391304348,
          4.565217391304348
        ],
        "common_position_by_pass": [
          89.13043478260869,
          89.13043478260869,
          89.13043478260869
        ],
        "display_mean": 4.565217391304348,
        "display_scale_position": 89.13043478260869,
        "repeat_statistics": {
          "n": 3,
          "mean": 89.13043478260869,
          "sample_sd": 0.0,
          "min": 89.13043478260869,
          "max": 89.13043478260869,
          "mean_minus_sd": 89.13043478260869
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 20,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": -0.5887681159420168
    },
    {
      "code": "C",
      "name": "Conscientiousness",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 24,
      "common_items": 24,
      "common_item_ids": [
        "ipip005",
        "ipip010",
        "ipip015",
        "ipip020",
        "ipip025",
        "ipip030",
        "ipip035",
        "ipip040",
        "ipip045",
        "ipip050",
        "ipip055",
        "ipip060",
        "ipip065",
        "ipip070",
        "ipip075",
        "ipip080",
        "ipip085",
        "ipip090",
        "ipip095",
        "ipip100",
        "ipip105",
        "ipip110",
        "ipip115",
        "ipip120"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 24,
          "missing_count": 0,
          "observed_sum": 100,
          "available_mean": 4.166666666666667,
          "available_scale_position": 79.16666666666667,
          "official_complete_raw_total": 100,
          "full_raw_total_bounds": [
            100,
            100
          ],
          "full_scale_position_bounds": [
            79.16666666666667,
            79.16666666666667
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 24,
          "missing_count": 0,
          "observed_sum": 99,
          "available_mean": 4.125,
          "available_scale_position": 78.125,
          "official_complete_raw_total": 99,
          "full_raw_total_bounds": [
            99,
            99
          ],
          "full_scale_position_bounds": [
            78.125,
            78.125
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 24,
          "missing_count": 0,
          "observed_sum": 99,
          "available_mean": 4.125,
          "available_scale_position": 78.125,
          "official_complete_raw_total": 99,
          "full_raw_total_bounds": [
            99,
            99
          ],
          "full_scale_position_bounds": [
            78.125,
            78.125
          ]
        }
      ],
      "common_mean_by_pass": [
        4.166666666666667,
        4.125,
        4.125
      ],
      "common_position_by_pass": [
        79.16666666666667,
        78.125,
        78.125
      ],
      "display_mean": 4.138888888888889,
      "display_scale_position": 78.47222222222223,
      "repeat_statistics": {
        "n": 3,
        "mean": 78.47222222222223,
        "sample_sd": 0.6014065304058629,
        "min": 78.125,
        "max": 79.16666666666667,
        "mean_minus_sd": 77.87081569181636
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 24,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 24,
        "common_items": 23,
        "common_item_ids": [
          "ipip005",
          "ipip010",
          "ipip015",
          "ipip020",
          "ipip025",
          "ipip030",
          "ipip035",
          "ipip045",
          "ipip050",
          "ipip055",
          "ipip060",
          "ipip065",
          "ipip070",
          "ipip075",
          "ipip080",
          "ipip085",
          "ipip090",
          "ipip095",
          "ipip100",
          "ipip105",
          "ipip110",
          "ipip115",
          "ipip120"
        ],
        "omitted_item_ids": [
          "ipip040"
        ],
        "complete_all_passes": false,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 24,
            "missing_count": 0,
            "observed_sum": 100,
            "available_mean": 4.166666666666667,
            "available_scale_position": 79.16666666666667,
            "official_complete_raw_total": 100,
            "full_raw_total_bounds": [
              100,
              100
            ],
            "full_scale_position_bounds": [
              79.16666666666667,
              79.16666666666667
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 24,
            "missing_count": 0,
            "observed_sum": 99,
            "available_mean": 4.125,
            "available_scale_position": 78.125,
            "official_complete_raw_total": 99,
            "full_raw_total_bounds": [
              99,
              99
            ],
            "full_scale_position_bounds": [
              78.125,
              78.125
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 23,
            "missing_count": 1,
            "observed_sum": 96,
            "available_mean": 4.173913043478261,
            "available_scale_position": 79.34782608695652,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              97,
              101
            ],
            "full_scale_position_bounds": [
              76.04166666666667,
              80.20833333333333
            ]
          }
        ],
        "common_mean_by_pass": [
          4.173913043478261,
          4.173913043478261,
          4.173913043478261
        ],
        "common_position_by_pass": [
          79.34782608695652,
          79.34782608695652,
          79.34782608695652
        ],
        "display_mean": 4.173913043478261,
        "display_scale_position": 79.34782608695652,
        "repeat_statistics": {
          "n": 3,
          "mean": 79.34782608695652,
          "sample_sd": 0.0,
          "min": 79.34782608695652,
          "max": 79.34782608695652,
          "mean_minus_sd": 79.34782608695652
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 20,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": -0.8756038647342876
    }
  ],
  "facets": [
    {
      "code": "N1",
      "name": "Anxiety",
      "domain": "N",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip001",
        "ipip031",
        "ipip061",
        "ipip091"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 6,
          "available_mean": 1.5,
          "available_scale_position": 12.5,
          "official_complete_raw_total": 6,
          "full_raw_total_bounds": [
            6,
            6
          ],
          "full_scale_position_bounds": [
            12.5,
            12.5
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 6,
          "available_mean": 1.5,
          "available_scale_position": 12.5,
          "official_complete_raw_total": 6,
          "full_raw_total_bounds": [
            6,
            6
          ],
          "full_scale_position_bounds": [
            12.5,
            12.5
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 6,
          "available_mean": 1.5,
          "available_scale_position": 12.5,
          "official_complete_raw_total": 6,
          "full_raw_total_bounds": [
            6,
            6
          ],
          "full_scale_position_bounds": [
            12.5,
            12.5
          ]
        }
      ],
      "common_mean_by_pass": [
        1.5,
        1.5,
        1.5
      ],
      "common_position_by_pass": [
        12.5,
        12.5,
        12.5
      ],
      "display_mean": 1.5,
      "display_scale_position": 12.5,
      "repeat_statistics": {
        "n": 3,
        "mean": 12.5,
        "sample_sd": 0.0,
        "min": 12.5,
        "max": 12.5,
        "mean_minus_sd": 12.5
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 3,
        "common_item_ids": [
          "ipip001",
          "ipip061",
          "ipip091"
        ],
        "omitted_item_ids": [
          "ipip031"
        ],
        "complete_all_passes": false,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 6,
            "available_mean": 1.5,
            "available_scale_position": 12.5,
            "official_complete_raw_total": 6,
            "full_raw_total_bounds": [
              6,
              6
            ],
            "full_scale_position_bounds": [
              12.5,
              12.5
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 6,
            "available_mean": 1.5,
            "available_scale_position": 12.5,
            "official_complete_raw_total": 6,
            "full_raw_total_bounds": [
              6,
              6
            ],
            "full_scale_position_bounds": [
              12.5,
              12.5
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 3,
            "missing_count": 1,
            "observed_sum": 5,
            "available_mean": 1.6666666666666667,
            "available_scale_position": 16.666666666666668,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              6,
              10
            ],
            "full_scale_position_bounds": [
              12.5,
              37.5
            ]
          }
        ],
        "common_mean_by_pass": [
          1.6666666666666667,
          1.6666666666666667,
          1.6666666666666667
        ],
        "common_position_by_pass": [
          16.666666666666668,
          16.666666666666668,
          16.666666666666668
        ],
        "display_mean": 1.6666666666666667,
        "display_scale_position": 16.666666666666668,
        "repeat_statistics": {
          "n": 3,
          "mean": 16.666666666666668,
          "sample_sd": 0.0,
          "min": 16.666666666666668,
          "max": 16.666666666666668,
          "mean_minus_sd": 16.666666666666668
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": -4.166666666666668
    },
    {
      "code": "N2",
      "name": "Anger",
      "domain": "N",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip006",
        "ipip036",
        "ipip066",
        "ipip096"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 7,
          "available_mean": 1.75,
          "available_scale_position": 18.75,
          "official_complete_raw_total": 7,
          "full_raw_total_bounds": [
            7,
            7
          ],
          "full_scale_position_bounds": [
            18.75,
            18.75
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 7,
          "available_mean": 1.75,
          "available_scale_position": 18.75,
          "official_complete_raw_total": 7,
          "full_raw_total_bounds": [
            7,
            7
          ],
          "full_scale_position_bounds": [
            18.75,
            18.75
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 7,
          "available_mean": 1.75,
          "available_scale_position": 18.75,
          "official_complete_raw_total": 7,
          "full_raw_total_bounds": [
            7,
            7
          ],
          "full_scale_position_bounds": [
            18.75,
            18.75
          ]
        }
      ],
      "common_mean_by_pass": [
        1.75,
        1.75,
        1.75
      ],
      "common_position_by_pass": [
        18.75,
        18.75,
        18.75
      ],
      "display_mean": 1.75,
      "display_scale_position": 18.75,
      "repeat_statistics": {
        "n": 3,
        "mean": 18.75,
        "sample_sd": 0.0,
        "min": 18.75,
        "max": 18.75,
        "mean_minus_sd": 18.75
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip006",
          "ipip036",
          "ipip066",
          "ipip096"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 7,
            "available_mean": 1.75,
            "available_scale_position": 18.75,
            "official_complete_raw_total": 7,
            "full_raw_total_bounds": [
              7,
              7
            ],
            "full_scale_position_bounds": [
              18.75,
              18.75
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 7,
            "available_mean": 1.75,
            "available_scale_position": 18.75,
            "official_complete_raw_total": 7,
            "full_raw_total_bounds": [
              7,
              7
            ],
            "full_scale_position_bounds": [
              18.75,
              18.75
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 7,
            "available_mean": 1.75,
            "available_scale_position": 18.75,
            "official_complete_raw_total": 7,
            "full_raw_total_bounds": [
              7,
              7
            ],
            "full_scale_position_bounds": [
              18.75,
              18.75
            ]
          }
        ],
        "common_mean_by_pass": [
          1.75,
          1.75,
          1.75
        ],
        "common_position_by_pass": [
          18.75,
          18.75,
          18.75
        ],
        "display_mean": 1.75,
        "display_scale_position": 18.75,
        "repeat_statistics": {
          "n": 3,
          "mean": 18.75,
          "sample_sd": 0.0,
          "min": 18.75,
          "max": 18.75,
          "mean_minus_sd": 18.75
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0.0
    },
    {
      "code": "N3",
      "name": "Depression",
      "domain": "N",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip011",
        "ipip041",
        "ipip071",
        "ipip101"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 5,
          "available_mean": 1.25,
          "available_scale_position": 6.25,
          "official_complete_raw_total": 5,
          "full_raw_total_bounds": [
            5,
            5
          ],
          "full_scale_position_bounds": [
            6.25,
            6.25
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 5,
          "available_mean": 1.25,
          "available_scale_position": 6.25,
          "official_complete_raw_total": 5,
          "full_raw_total_bounds": [
            5,
            5
          ],
          "full_scale_position_bounds": [
            6.25,
            6.25
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 5,
          "available_mean": 1.25,
          "available_scale_position": 6.25,
          "official_complete_raw_total": 5,
          "full_raw_total_bounds": [
            5,
            5
          ],
          "full_scale_position_bounds": [
            6.25,
            6.25
          ]
        }
      ],
      "common_mean_by_pass": [
        1.25,
        1.25,
        1.25
      ],
      "common_position_by_pass": [
        6.25,
        6.25,
        6.25
      ],
      "display_mean": 1.25,
      "display_scale_position": 6.25,
      "repeat_statistics": {
        "n": 3,
        "mean": 6.25,
        "sample_sd": 0.0,
        "min": 6.25,
        "max": 6.25,
        "mean_minus_sd": 6.25
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": false,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 2,
        "common_item_ids": [
          "ipip011",
          "ipip101"
        ],
        "omitted_item_ids": [
          "ipip041",
          "ipip071"
        ],
        "complete_all_passes": false,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 3,
            "missing_count": 1,
            "observed_sum": 4,
            "available_mean": 1.3333333333333333,
            "available_scale_position": 8.333333333333332,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              5,
              9
            ],
            "full_scale_position_bounds": [
              6.25,
              31.25
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 3,
            "missing_count": 1,
            "observed_sum": 4,
            "available_mean": 1.3333333333333333,
            "available_scale_position": 8.333333333333332,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              5,
              9
            ],
            "full_scale_position_bounds": [
              6.25,
              31.25
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 5,
            "available_mean": 1.25,
            "available_scale_position": 6.25,
            "official_complete_raw_total": 5,
            "full_raw_total_bounds": [
              5,
              5
            ],
            "full_scale_position_bounds": [
              6.25,
              6.25
            ]
          }
        ],
        "common_mean_by_pass": [
          1.5,
          1.5,
          1.5
        ],
        "common_position_by_pass": [
          12.5,
          12.5,
          12.5
        ],
        "display_mean": null,
        "display_scale_position": null,
        "repeat_statistics": {
          "n": 0,
          "mean": null,
          "sample_sd": null,
          "min": null,
          "max": null,
          "mean_minus_sd": null
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": null
    },
    {
      "code": "N4",
      "name": "Self-Consciousness",
      "domain": "N",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip016",
        "ipip046",
        "ipip076",
        "ipip106"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 5,
          "available_mean": 1.25,
          "available_scale_position": 6.25,
          "official_complete_raw_total": 5,
          "full_raw_total_bounds": [
            5,
            5
          ],
          "full_scale_position_bounds": [
            6.25,
            6.25
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 5,
          "available_mean": 1.25,
          "available_scale_position": 6.25,
          "official_complete_raw_total": 5,
          "full_raw_total_bounds": [
            5,
            5
          ],
          "full_scale_position_bounds": [
            6.25,
            6.25
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 5,
          "available_mean": 1.25,
          "available_scale_position": 6.25,
          "official_complete_raw_total": 5,
          "full_raw_total_bounds": [
            5,
            5
          ],
          "full_scale_position_bounds": [
            6.25,
            6.25
          ]
        }
      ],
      "common_mean_by_pass": [
        1.25,
        1.25,
        1.25
      ],
      "common_position_by_pass": [
        6.25,
        6.25,
        6.25
      ],
      "display_mean": 1.25,
      "display_scale_position": 6.25,
      "repeat_statistics": {
        "n": 3,
        "mean": 6.25,
        "sample_sd": 0.0,
        "min": 6.25,
        "max": 6.25,
        "mean_minus_sd": 6.25
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip016",
          "ipip046",
          "ipip076",
          "ipip106"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 5,
            "available_mean": 1.25,
            "available_scale_position": 6.25,
            "official_complete_raw_total": 5,
            "full_raw_total_bounds": [
              5,
              5
            ],
            "full_scale_position_bounds": [
              6.25,
              6.25
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 5,
            "available_mean": 1.25,
            "available_scale_position": 6.25,
            "official_complete_raw_total": 5,
            "full_raw_total_bounds": [
              5,
              5
            ],
            "full_scale_position_bounds": [
              6.25,
              6.25
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 5,
            "available_mean": 1.25,
            "available_scale_position": 6.25,
            "official_complete_raw_total": 5,
            "full_raw_total_bounds": [
              5,
              5
            ],
            "full_scale_position_bounds": [
              6.25,
              6.25
            ]
          }
        ],
        "common_mean_by_pass": [
          1.25,
          1.25,
          1.25
        ],
        "common_position_by_pass": [
          6.25,
          6.25,
          6.25
        ],
        "display_mean": 1.25,
        "display_scale_position": 6.25,
        "repeat_statistics": {
          "n": 3,
          "mean": 6.25,
          "sample_sd": 0.0,
          "min": 6.25,
          "max": 6.25,
          "mean_minus_sd": 6.25
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0.0
    },
    {
      "code": "N5",
      "name": "Immoderation",
      "domain": "N",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip021",
        "ipip051",
        "ipip081",
        "ipip111"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 9,
          "available_mean": 2.25,
          "available_scale_position": 31.25,
          "official_complete_raw_total": 9,
          "full_raw_total_bounds": [
            9,
            9
          ],
          "full_scale_position_bounds": [
            31.25,
            31.25
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 9,
          "available_mean": 2.25,
          "available_scale_position": 31.25,
          "official_complete_raw_total": 9,
          "full_raw_total_bounds": [
            9,
            9
          ],
          "full_scale_position_bounds": [
            31.25,
            31.25
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 9,
          "available_mean": 2.25,
          "available_scale_position": 31.25,
          "official_complete_raw_total": 9,
          "full_raw_total_bounds": [
            9,
            9
          ],
          "full_scale_position_bounds": [
            31.25,
            31.25
          ]
        }
      ],
      "common_mean_by_pass": [
        2.25,
        2.25,
        2.25
      ],
      "common_position_by_pass": [
        31.25,
        31.25,
        31.25
      ],
      "display_mean": 2.25,
      "display_scale_position": 31.25,
      "repeat_statistics": {
        "n": 3,
        "mean": 31.25,
        "sample_sd": 0.0,
        "min": 31.25,
        "max": 31.25,
        "mean_minus_sd": 31.25
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 3,
        "common_item_ids": [
          "ipip051",
          "ipip081",
          "ipip111"
        ],
        "omitted_item_ids": [
          "ipip021"
        ],
        "complete_all_passes": false,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 9,
            "available_mean": 2.25,
            "available_scale_position": 31.25,
            "official_complete_raw_total": 9,
            "full_raw_total_bounds": [
              9,
              9
            ],
            "full_scale_position_bounds": [
              31.25,
              31.25
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 3,
            "missing_count": 1,
            "observed_sum": 8,
            "available_mean": 2.6666666666666665,
            "available_scale_position": 41.666666666666664,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              9,
              13
            ],
            "full_scale_position_bounds": [
              31.25,
              56.25
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 9,
            "available_mean": 2.25,
            "available_scale_position": 31.25,
            "official_complete_raw_total": 9,
            "full_raw_total_bounds": [
              9,
              9
            ],
            "full_scale_position_bounds": [
              31.25,
              31.25
            ]
          }
        ],
        "common_mean_by_pass": [
          2.6666666666666665,
          2.6666666666666665,
          2.6666666666666665
        ],
        "common_position_by_pass": [
          41.666666666666664,
          41.666666666666664,
          41.666666666666664
        ],
        "display_mean": 2.6666666666666665,
        "display_scale_position": 41.666666666666664,
        "repeat_statistics": {
          "n": 3,
          "mean": 41.666666666666664,
          "sample_sd": 0.0,
          "min": 41.666666666666664,
          "max": 41.666666666666664,
          "mean_minus_sd": 41.666666666666664
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": -10.416666666666664
    },
    {
      "code": "N6",
      "name": "Vulnerability",
      "domain": "N",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip026",
        "ipip056",
        "ipip086",
        "ipip116"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 6,
          "available_mean": 1.5,
          "available_scale_position": 12.5,
          "official_complete_raw_total": 6,
          "full_raw_total_bounds": [
            6,
            6
          ],
          "full_scale_position_bounds": [
            12.5,
            12.5
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 6,
          "available_mean": 1.5,
          "available_scale_position": 12.5,
          "official_complete_raw_total": 6,
          "full_raw_total_bounds": [
            6,
            6
          ],
          "full_scale_position_bounds": [
            12.5,
            12.5
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 6,
          "available_mean": 1.5,
          "available_scale_position": 12.5,
          "official_complete_raw_total": 6,
          "full_raw_total_bounds": [
            6,
            6
          ],
          "full_scale_position_bounds": [
            12.5,
            12.5
          ]
        }
      ],
      "common_mean_by_pass": [
        1.5,
        1.5,
        1.5
      ],
      "common_position_by_pass": [
        12.5,
        12.5,
        12.5
      ],
      "display_mean": 1.5,
      "display_scale_position": 12.5,
      "repeat_statistics": {
        "n": 3,
        "mean": 12.5,
        "sample_sd": 0.0,
        "min": 12.5,
        "max": 12.5,
        "mean_minus_sd": 12.5
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip026",
          "ipip056",
          "ipip086",
          "ipip116"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 6,
            "available_mean": 1.5,
            "available_scale_position": 12.5,
            "official_complete_raw_total": 6,
            "full_raw_total_bounds": [
              6,
              6
            ],
            "full_scale_position_bounds": [
              12.5,
              12.5
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 6,
            "available_mean": 1.5,
            "available_scale_position": 12.5,
            "official_complete_raw_total": 6,
            "full_raw_total_bounds": [
              6,
              6
            ],
            "full_scale_position_bounds": [
              12.5,
              12.5
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 6,
            "available_mean": 1.5,
            "available_scale_position": 12.5,
            "official_complete_raw_total": 6,
            "full_raw_total_bounds": [
              6,
              6
            ],
            "full_scale_position_bounds": [
              12.5,
              12.5
            ]
          }
        ],
        "common_mean_by_pass": [
          1.5,
          1.5,
          1.5
        ],
        "common_position_by_pass": [
          12.5,
          12.5,
          12.5
        ],
        "display_mean": 1.5,
        "display_scale_position": 12.5,
        "repeat_statistics": {
          "n": 3,
          "mean": 12.5,
          "sample_sd": 0.0,
          "min": 12.5,
          "max": 12.5,
          "mean_minus_sd": 12.5
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0.0
    },
    {
      "code": "E1",
      "name": "Friendliness",
      "domain": "E",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip002",
        "ipip032",
        "ipip062",
        "ipip092"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        }
      ],
      "common_mean_by_pass": [
        4.5,
        4.5,
        4.5
      ],
      "common_position_by_pass": [
        87.5,
        87.5,
        87.5
      ],
      "display_mean": 4.5,
      "display_scale_position": 87.5,
      "repeat_statistics": {
        "n": 3,
        "mean": 87.5,
        "sample_sd": 0.0,
        "min": 87.5,
        "max": 87.5,
        "mean_minus_sd": 87.5
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip002",
          "ipip032",
          "ipip062",
          "ipip092"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          }
        ],
        "common_mean_by_pass": [
          4.5,
          4.5,
          4.5
        ],
        "common_position_by_pass": [
          87.5,
          87.5,
          87.5
        ],
        "display_mean": 4.5,
        "display_scale_position": 87.5,
        "repeat_statistics": {
          "n": 3,
          "mean": 87.5,
          "sample_sd": 0.0,
          "min": 87.5,
          "max": 87.5,
          "mean_minus_sd": 87.5
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0.0
    },
    {
      "code": "E2",
      "name": "Gregariousness",
      "domain": "E",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip007",
        "ipip037",
        "ipip067",
        "ipip097"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 13,
          "available_mean": 3.25,
          "available_scale_position": 56.25,
          "official_complete_raw_total": 13,
          "full_raw_total_bounds": [
            13,
            13
          ],
          "full_scale_position_bounds": [
            56.25,
            56.25
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 13,
          "available_mean": 3.25,
          "available_scale_position": 56.25,
          "official_complete_raw_total": 13,
          "full_raw_total_bounds": [
            13,
            13
          ],
          "full_scale_position_bounds": [
            56.25,
            56.25
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 13,
          "available_mean": 3.25,
          "available_scale_position": 56.25,
          "official_complete_raw_total": 13,
          "full_raw_total_bounds": [
            13,
            13
          ],
          "full_scale_position_bounds": [
            56.25,
            56.25
          ]
        }
      ],
      "common_mean_by_pass": [
        3.25,
        3.25,
        3.25
      ],
      "common_position_by_pass": [
        56.25,
        56.25,
        56.25
      ],
      "display_mean": 3.25,
      "display_scale_position": 56.25,
      "repeat_statistics": {
        "n": 3,
        "mean": 56.25,
        "sample_sd": 0.0,
        "min": 56.25,
        "max": 56.25,
        "mean_minus_sd": 56.25
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 3,
        "common_item_ids": [
          "ipip007",
          "ipip067",
          "ipip097"
        ],
        "omitted_item_ids": [
          "ipip037"
        ],
        "complete_all_passes": false,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 13,
            "available_mean": 3.25,
            "available_scale_position": 56.25,
            "official_complete_raw_total": 13,
            "full_raw_total_bounds": [
              13,
              13
            ],
            "full_scale_position_bounds": [
              56.25,
              56.25
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 3,
            "missing_count": 1,
            "observed_sum": 11,
            "available_mean": 3.6666666666666665,
            "available_scale_position": 66.66666666666666,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              12,
              16
            ],
            "full_scale_position_bounds": [
              50.0,
              75.0
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 13,
            "available_mean": 3.25,
            "available_scale_position": 56.25,
            "official_complete_raw_total": 13,
            "full_raw_total_bounds": [
              13,
              13
            ],
            "full_scale_position_bounds": [
              56.25,
              56.25
            ]
          }
        ],
        "common_mean_by_pass": [
          3.6666666666666665,
          3.6666666666666665,
          3.6666666666666665
        ],
        "common_position_by_pass": [
          66.66666666666666,
          66.66666666666666,
          66.66666666666666
        ],
        "display_mean": 3.6666666666666665,
        "display_scale_position": 66.66666666666666,
        "repeat_statistics": {
          "n": 3,
          "mean": 66.66666666666666,
          "sample_sd": 0.0,
          "min": 66.66666666666666,
          "max": 66.66666666666666,
          "mean_minus_sd": 66.66666666666666
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": -10.416666666666657
    },
    {
      "code": "E3",
      "name": "Assertiveness",
      "domain": "E",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip012",
        "ipip042",
        "ipip072",
        "ipip102"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 11,
          "available_mean": 2.75,
          "available_scale_position": 43.75,
          "official_complete_raw_total": 11,
          "full_raw_total_bounds": [
            11,
            11
          ],
          "full_scale_position_bounds": [
            43.75,
            43.75
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 11,
          "available_mean": 2.75,
          "available_scale_position": 43.75,
          "official_complete_raw_total": 11,
          "full_raw_total_bounds": [
            11,
            11
          ],
          "full_scale_position_bounds": [
            43.75,
            43.75
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 11,
          "available_mean": 2.75,
          "available_scale_position": 43.75,
          "official_complete_raw_total": 11,
          "full_raw_total_bounds": [
            11,
            11
          ],
          "full_scale_position_bounds": [
            43.75,
            43.75
          ]
        }
      ],
      "common_mean_by_pass": [
        2.75,
        2.75,
        2.75
      ],
      "common_position_by_pass": [
        43.75,
        43.75,
        43.75
      ],
      "display_mean": 2.75,
      "display_scale_position": 43.75,
      "repeat_statistics": {
        "n": 3,
        "mean": 43.75,
        "sample_sd": 0.0,
        "min": 43.75,
        "max": 43.75,
        "mean_minus_sd": 43.75
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip012",
          "ipip042",
          "ipip072",
          "ipip102"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 11,
            "available_mean": 2.75,
            "available_scale_position": 43.75,
            "official_complete_raw_total": 11,
            "full_raw_total_bounds": [
              11,
              11
            ],
            "full_scale_position_bounds": [
              43.75,
              43.75
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 11,
            "available_mean": 2.75,
            "available_scale_position": 43.75,
            "official_complete_raw_total": 11,
            "full_raw_total_bounds": [
              11,
              11
            ],
            "full_scale_position_bounds": [
              43.75,
              43.75
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 11,
            "available_mean": 2.75,
            "available_scale_position": 43.75,
            "official_complete_raw_total": 11,
            "full_raw_total_bounds": [
              11,
              11
            ],
            "full_scale_position_bounds": [
              43.75,
              43.75
            ]
          }
        ],
        "common_mean_by_pass": [
          2.75,
          2.75,
          2.75
        ],
        "common_position_by_pass": [
          43.75,
          43.75,
          43.75
        ],
        "display_mean": 2.75,
        "display_scale_position": 43.75,
        "repeat_statistics": {
          "n": 3,
          "mean": 43.75,
          "sample_sd": 0.0,
          "min": 43.75,
          "max": 43.75,
          "mean_minus_sd": 43.75
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0.0
    },
    {
      "code": "E4",
      "name": "Activity Level",
      "domain": "E",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip017",
        "ipip047",
        "ipip077",
        "ipip107"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 16,
          "available_mean": 4,
          "available_scale_position": 75,
          "official_complete_raw_total": 16,
          "full_raw_total_bounds": [
            16,
            16
          ],
          "full_scale_position_bounds": [
            75.0,
            75.0
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 16,
          "available_mean": 4,
          "available_scale_position": 75,
          "official_complete_raw_total": 16,
          "full_raw_total_bounds": [
            16,
            16
          ],
          "full_scale_position_bounds": [
            75.0,
            75.0
          ]
        }
      ],
      "common_mean_by_pass": [
        4,
        4.5,
        4
      ],
      "common_position_by_pass": [
        75,
        87.5,
        75
      ],
      "display_mean": 4.166666666666667,
      "display_scale_position": 79.16666666666667,
      "repeat_statistics": {
        "n": 3,
        "mean": 79.16666666666667,
        "sample_sd": 7.216878364870322,
        "min": 75,
        "max": 87.5,
        "mean_minus_sd": 71.94978830179635
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 3,
        "common_item_ids": [
          "ipip017",
          "ipip077",
          "ipip107"
        ],
        "omitted_item_ids": [
          "ipip047"
        ],
        "complete_all_passes": false,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 3,
            "missing_count": 1,
            "observed_sum": 12,
            "available_mean": 4,
            "available_scale_position": 75,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              13,
              17
            ],
            "full_scale_position_bounds": [
              56.25,
              81.25
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 16,
            "available_mean": 4,
            "available_scale_position": 75,
            "official_complete_raw_total": 16,
            "full_raw_total_bounds": [
              16,
              16
            ],
            "full_scale_position_bounds": [
              75.0,
              75.0
            ]
          }
        ],
        "common_mean_by_pass": [
          4,
          4.666666666666667,
          4
        ],
        "common_position_by_pass": [
          75,
          91.66666666666667,
          75
        ],
        "display_mean": 4.222222222222222,
        "display_scale_position": 80.55555555555556,
        "repeat_statistics": {
          "n": 3,
          "mean": 80.55555555555556,
          "sample_sd": 9.622504486493765,
          "min": 75,
          "max": 91.66666666666667,
          "mean_minus_sd": 70.93305106906179
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": -1.3888888888888857
    },
    {
      "code": "E5",
      "name": "Excitement-Seeking",
      "domain": "E",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip022",
        "ipip052",
        "ipip082",
        "ipip112"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 13,
          "available_mean": 3.25,
          "available_scale_position": 56.25,
          "official_complete_raw_total": 13,
          "full_raw_total_bounds": [
            13,
            13
          ],
          "full_scale_position_bounds": [
            56.25,
            56.25
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 13,
          "available_mean": 3.25,
          "available_scale_position": 56.25,
          "official_complete_raw_total": 13,
          "full_raw_total_bounds": [
            13,
            13
          ],
          "full_scale_position_bounds": [
            56.25,
            56.25
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 13,
          "available_mean": 3.25,
          "available_scale_position": 56.25,
          "official_complete_raw_total": 13,
          "full_raw_total_bounds": [
            13,
            13
          ],
          "full_scale_position_bounds": [
            56.25,
            56.25
          ]
        }
      ],
      "common_mean_by_pass": [
        3.25,
        3.25,
        3.25
      ],
      "common_position_by_pass": [
        56.25,
        56.25,
        56.25
      ],
      "display_mean": 3.25,
      "display_scale_position": 56.25,
      "repeat_statistics": {
        "n": 3,
        "mean": 56.25,
        "sample_sd": 0.0,
        "min": 56.25,
        "max": 56.25,
        "mean_minus_sd": 56.25
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip022",
          "ipip052",
          "ipip082",
          "ipip112"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 13,
            "available_mean": 3.25,
            "available_scale_position": 56.25,
            "official_complete_raw_total": 13,
            "full_raw_total_bounds": [
              13,
              13
            ],
            "full_scale_position_bounds": [
              56.25,
              56.25
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 13,
            "available_mean": 3.25,
            "available_scale_position": 56.25,
            "official_complete_raw_total": 13,
            "full_raw_total_bounds": [
              13,
              13
            ],
            "full_scale_position_bounds": [
              56.25,
              56.25
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 13,
            "available_mean": 3.25,
            "available_scale_position": 56.25,
            "official_complete_raw_total": 13,
            "full_raw_total_bounds": [
              13,
              13
            ],
            "full_scale_position_bounds": [
              56.25,
              56.25
            ]
          }
        ],
        "common_mean_by_pass": [
          3.25,
          3.25,
          3.25
        ],
        "common_position_by_pass": [
          56.25,
          56.25,
          56.25
        ],
        "display_mean": 3.25,
        "display_scale_position": 56.25,
        "repeat_statistics": {
          "n": 3,
          "mean": 56.25,
          "sample_sd": 0.0,
          "min": 56.25,
          "max": 56.25,
          "mean_minus_sd": 56.25
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0.0
    },
    {
      "code": "E6",
      "name": "Cheerfulness",
      "domain": "E",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip027",
        "ipip057",
        "ipip087",
        "ipip117"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 13,
          "available_mean": 3.25,
          "available_scale_position": 56.25,
          "official_complete_raw_total": 13,
          "full_raw_total_bounds": [
            13,
            13
          ],
          "full_scale_position_bounds": [
            56.25,
            56.25
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 13,
          "available_mean": 3.25,
          "available_scale_position": 56.25,
          "official_complete_raw_total": 13,
          "full_raw_total_bounds": [
            13,
            13
          ],
          "full_scale_position_bounds": [
            56.25,
            56.25
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 13,
          "available_mean": 3.25,
          "available_scale_position": 56.25,
          "official_complete_raw_total": 13,
          "full_raw_total_bounds": [
            13,
            13
          ],
          "full_scale_position_bounds": [
            56.25,
            56.25
          ]
        }
      ],
      "common_mean_by_pass": [
        3.25,
        3.25,
        3.25
      ],
      "common_position_by_pass": [
        56.25,
        56.25,
        56.25
      ],
      "display_mean": 3.25,
      "display_scale_position": 56.25,
      "repeat_statistics": {
        "n": 3,
        "mean": 56.25,
        "sample_sd": 0.0,
        "min": 56.25,
        "max": 56.25,
        "mean_minus_sd": 56.25
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 3,
        "common_item_ids": [
          "ipip027",
          "ipip057",
          "ipip087"
        ],
        "omitted_item_ids": [
          "ipip117"
        ],
        "complete_all_passes": false,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 3,
            "missing_count": 1,
            "observed_sum": 9,
            "available_mean": 3,
            "available_scale_position": 50,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              10,
              14
            ],
            "full_scale_position_bounds": [
              37.5,
              62.5
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 3,
            "missing_count": 1,
            "observed_sum": 9,
            "available_mean": 3,
            "available_scale_position": 50,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              10,
              14
            ],
            "full_scale_position_bounds": [
              37.5,
              62.5
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 13,
            "available_mean": 3.25,
            "available_scale_position": 56.25,
            "official_complete_raw_total": 13,
            "full_raw_total_bounds": [
              13,
              13
            ],
            "full_scale_position_bounds": [
              56.25,
              56.25
            ]
          }
        ],
        "common_mean_by_pass": [
          3,
          3,
          3
        ],
        "common_position_by_pass": [
          50,
          50,
          50
        ],
        "display_mean": 3,
        "display_scale_position": 50,
        "repeat_statistics": {
          "n": 3,
          "mean": 50,
          "sample_sd": 0.0,
          "min": 50,
          "max": 50,
          "mean_minus_sd": 50.0
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 6.25
    },
    {
      "code": "O1",
      "name": "Imagination",
      "domain": "O",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip003",
        "ipip033",
        "ipip063",
        "ipip093"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 15,
          "available_mean": 3.75,
          "available_scale_position": 68.75,
          "official_complete_raw_total": 15,
          "full_raw_total_bounds": [
            15,
            15
          ],
          "full_scale_position_bounds": [
            68.75,
            68.75
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 15,
          "available_mean": 3.75,
          "available_scale_position": 68.75,
          "official_complete_raw_total": 15,
          "full_raw_total_bounds": [
            15,
            15
          ],
          "full_scale_position_bounds": [
            68.75,
            68.75
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 15,
          "available_mean": 3.75,
          "available_scale_position": 68.75,
          "official_complete_raw_total": 15,
          "full_raw_total_bounds": [
            15,
            15
          ],
          "full_scale_position_bounds": [
            68.75,
            68.75
          ]
        }
      ],
      "common_mean_by_pass": [
        3.75,
        3.75,
        3.75
      ],
      "common_position_by_pass": [
        68.75,
        68.75,
        68.75
      ],
      "display_mean": 3.75,
      "display_scale_position": 68.75,
      "repeat_statistics": {
        "n": 3,
        "mean": 68.75,
        "sample_sd": 0.0,
        "min": 68.75,
        "max": 68.75,
        "mean_minus_sd": 68.75
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip003",
          "ipip033",
          "ipip063",
          "ipip093"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 15,
            "available_mean": 3.75,
            "available_scale_position": 68.75,
            "official_complete_raw_total": 15,
            "full_raw_total_bounds": [
              15,
              15
            ],
            "full_scale_position_bounds": [
              68.75,
              68.75
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 15,
            "available_mean": 3.75,
            "available_scale_position": 68.75,
            "official_complete_raw_total": 15,
            "full_raw_total_bounds": [
              15,
              15
            ],
            "full_scale_position_bounds": [
              68.75,
              68.75
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 15,
            "available_mean": 3.75,
            "available_scale_position": 68.75,
            "official_complete_raw_total": 15,
            "full_raw_total_bounds": [
              15,
              15
            ],
            "full_scale_position_bounds": [
              68.75,
              68.75
            ]
          }
        ],
        "common_mean_by_pass": [
          3.75,
          3.75,
          3.75
        ],
        "common_position_by_pass": [
          68.75,
          68.75,
          68.75
        ],
        "display_mean": 3.75,
        "display_scale_position": 68.75,
        "repeat_statistics": {
          "n": 3,
          "mean": 68.75,
          "sample_sd": 0.0,
          "min": 68.75,
          "max": 68.75,
          "mean_minus_sd": 68.75
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0.0
    },
    {
      "code": "O2",
      "name": "Artistic Interests",
      "domain": "O",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip008",
        "ipip038",
        "ipip068",
        "ipip098"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        }
      ],
      "common_mean_by_pass": [
        4.5,
        4.5,
        4.5
      ],
      "common_position_by_pass": [
        87.5,
        87.5,
        87.5
      ],
      "display_mean": 4.5,
      "display_scale_position": 87.5,
      "repeat_statistics": {
        "n": 3,
        "mean": 87.5,
        "sample_sd": 0.0,
        "min": 87.5,
        "max": 87.5,
        "mean_minus_sd": 87.5
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip008",
          "ipip038",
          "ipip068",
          "ipip098"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          }
        ],
        "common_mean_by_pass": [
          4.5,
          4.5,
          4.5
        ],
        "common_position_by_pass": [
          87.5,
          87.5,
          87.5
        ],
        "display_mean": 4.5,
        "display_scale_position": 87.5,
        "repeat_statistics": {
          "n": 3,
          "mean": 87.5,
          "sample_sd": 0.0,
          "min": 87.5,
          "max": 87.5,
          "mean_minus_sd": 87.5
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0.0
    },
    {
      "code": "O3",
      "name": "Emotionality",
      "domain": "O",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip013",
        "ipip043",
        "ipip073",
        "ipip103"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 16,
          "available_mean": 4,
          "available_scale_position": 75,
          "official_complete_raw_total": 16,
          "full_raw_total_bounds": [
            16,
            16
          ],
          "full_scale_position_bounds": [
            75.0,
            75.0
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 16,
          "available_mean": 4,
          "available_scale_position": 75,
          "official_complete_raw_total": 16,
          "full_raw_total_bounds": [
            16,
            16
          ],
          "full_scale_position_bounds": [
            75.0,
            75.0
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 16,
          "available_mean": 4,
          "available_scale_position": 75,
          "official_complete_raw_total": 16,
          "full_raw_total_bounds": [
            16,
            16
          ],
          "full_scale_position_bounds": [
            75.0,
            75.0
          ]
        }
      ],
      "common_mean_by_pass": [
        4,
        4,
        4
      ],
      "common_position_by_pass": [
        75,
        75,
        75
      ],
      "display_mean": 4,
      "display_scale_position": 75,
      "repeat_statistics": {
        "n": 3,
        "mean": 75,
        "sample_sd": 0.0,
        "min": 75,
        "max": 75,
        "mean_minus_sd": 75.0
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip013",
          "ipip043",
          "ipip073",
          "ipip103"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 16,
            "available_mean": 4,
            "available_scale_position": 75,
            "official_complete_raw_total": 16,
            "full_raw_total_bounds": [
              16,
              16
            ],
            "full_scale_position_bounds": [
              75.0,
              75.0
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 16,
            "available_mean": 4,
            "available_scale_position": 75,
            "official_complete_raw_total": 16,
            "full_raw_total_bounds": [
              16,
              16
            ],
            "full_scale_position_bounds": [
              75.0,
              75.0
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 16,
            "available_mean": 4,
            "available_scale_position": 75,
            "official_complete_raw_total": 16,
            "full_raw_total_bounds": [
              16,
              16
            ],
            "full_scale_position_bounds": [
              75.0,
              75.0
            ]
          }
        ],
        "common_mean_by_pass": [
          4,
          4,
          4
        ],
        "common_position_by_pass": [
          75,
          75,
          75
        ],
        "display_mean": 4,
        "display_scale_position": 75,
        "repeat_statistics": {
          "n": 3,
          "mean": 75,
          "sample_sd": 0.0,
          "min": 75,
          "max": 75,
          "mean_minus_sd": 75.0
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0
    },
    {
      "code": "O4",
      "name": "Adventurousness",
      "domain": "O",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip018",
        "ipip048",
        "ipip078",
        "ipip108"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 20,
          "available_mean": 5,
          "available_scale_position": 100,
          "official_complete_raw_total": 20,
          "full_raw_total_bounds": [
            20,
            20
          ],
          "full_scale_position_bounds": [
            100.0,
            100.0
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 19,
          "available_mean": 4.75,
          "available_scale_position": 93.75,
          "official_complete_raw_total": 19,
          "full_raw_total_bounds": [
            19,
            19
          ],
          "full_scale_position_bounds": [
            93.75,
            93.75
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 19,
          "available_mean": 4.75,
          "available_scale_position": 93.75,
          "official_complete_raw_total": 19,
          "full_raw_total_bounds": [
            19,
            19
          ],
          "full_scale_position_bounds": [
            93.75,
            93.75
          ]
        }
      ],
      "common_mean_by_pass": [
        5,
        4.75,
        4.75
      ],
      "common_position_by_pass": [
        100,
        93.75,
        93.75
      ],
      "display_mean": 4.833333333333333,
      "display_scale_position": 95.83333333333333,
      "repeat_statistics": {
        "n": 3,
        "mean": 95.83333333333333,
        "sample_sd": 3.608439182435161,
        "min": 93.75,
        "max": 100,
        "mean_minus_sd": 92.22489415089817
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 3,
        "common_item_ids": [
          "ipip018",
          "ipip048",
          "ipip078"
        ],
        "omitted_item_ids": [
          "ipip108"
        ],
        "complete_all_passes": false,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 20,
            "available_mean": 5,
            "available_scale_position": 100,
            "official_complete_raw_total": 20,
            "full_raw_total_bounds": [
              20,
              20
            ],
            "full_scale_position_bounds": [
              100.0,
              100.0
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 3,
            "missing_count": 1,
            "observed_sum": 15,
            "available_mean": 5,
            "available_scale_position": 100,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              16,
              20
            ],
            "full_scale_position_bounds": [
              75.0,
              100.0
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 19,
            "available_mean": 4.75,
            "available_scale_position": 93.75,
            "official_complete_raw_total": 19,
            "full_raw_total_bounds": [
              19,
              19
            ],
            "full_scale_position_bounds": [
              93.75,
              93.75
            ]
          }
        ],
        "common_mean_by_pass": [
          5,
          5,
          5
        ],
        "common_position_by_pass": [
          100,
          100,
          100
        ],
        "display_mean": 5,
        "display_scale_position": 100,
        "repeat_statistics": {
          "n": 3,
          "mean": 100,
          "sample_sd": 0.0,
          "min": 100,
          "max": 100,
          "mean_minus_sd": 100.0
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": -4.166666666666671
    },
    {
      "code": "O5",
      "name": "Intellect",
      "domain": "O",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip023",
        "ipip053",
        "ipip083",
        "ipip113"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 19,
          "available_mean": 4.75,
          "available_scale_position": 93.75,
          "official_complete_raw_total": 19,
          "full_raw_total_bounds": [
            19,
            19
          ],
          "full_scale_position_bounds": [
            93.75,
            93.75
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 19,
          "available_mean": 4.75,
          "available_scale_position": 93.75,
          "official_complete_raw_total": 19,
          "full_raw_total_bounds": [
            19,
            19
          ],
          "full_scale_position_bounds": [
            93.75,
            93.75
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 19,
          "available_mean": 4.75,
          "available_scale_position": 93.75,
          "official_complete_raw_total": 19,
          "full_raw_total_bounds": [
            19,
            19
          ],
          "full_scale_position_bounds": [
            93.75,
            93.75
          ]
        }
      ],
      "common_mean_by_pass": [
        4.75,
        4.75,
        4.75
      ],
      "common_position_by_pass": [
        93.75,
        93.75,
        93.75
      ],
      "display_mean": 4.75,
      "display_scale_position": 93.75,
      "repeat_statistics": {
        "n": 3,
        "mean": 93.75,
        "sample_sd": 0.0,
        "min": 93.75,
        "max": 93.75,
        "mean_minus_sd": 93.75
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip023",
          "ipip053",
          "ipip083",
          "ipip113"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 19,
            "available_mean": 4.75,
            "available_scale_position": 93.75,
            "official_complete_raw_total": 19,
            "full_raw_total_bounds": [
              19,
              19
            ],
            "full_scale_position_bounds": [
              93.75,
              93.75
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 19,
            "available_mean": 4.75,
            "available_scale_position": 93.75,
            "official_complete_raw_total": 19,
            "full_raw_total_bounds": [
              19,
              19
            ],
            "full_scale_position_bounds": [
              93.75,
              93.75
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 19,
            "available_mean": 4.75,
            "available_scale_position": 93.75,
            "official_complete_raw_total": 19,
            "full_raw_total_bounds": [
              19,
              19
            ],
            "full_scale_position_bounds": [
              93.75,
              93.75
            ]
          }
        ],
        "common_mean_by_pass": [
          4.75,
          4.75,
          4.75
        ],
        "common_position_by_pass": [
          93.75,
          93.75,
          93.75
        ],
        "display_mean": 4.75,
        "display_scale_position": 93.75,
        "repeat_statistics": {
          "n": 3,
          "mean": 93.75,
          "sample_sd": 0.0,
          "min": 93.75,
          "max": 93.75,
          "mean_minus_sd": 93.75
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0.0
    },
    {
      "code": "O6",
      "name": "Liberalism",
      "domain": "O",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip028",
        "ipip058",
        "ipip088",
        "ipip118"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 15,
          "available_mean": 3.75,
          "available_scale_position": 68.75,
          "official_complete_raw_total": 15,
          "full_raw_total_bounds": [
            15,
            15
          ],
          "full_scale_position_bounds": [
            68.75,
            68.75
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 15,
          "available_mean": 3.75,
          "available_scale_position": 68.75,
          "official_complete_raw_total": 15,
          "full_raw_total_bounds": [
            15,
            15
          ],
          "full_scale_position_bounds": [
            68.75,
            68.75
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 14,
          "available_mean": 3.5,
          "available_scale_position": 62.5,
          "official_complete_raw_total": 14,
          "full_raw_total_bounds": [
            14,
            14
          ],
          "full_scale_position_bounds": [
            62.5,
            62.5
          ]
        }
      ],
      "common_mean_by_pass": [
        3.75,
        3.75,
        3.5
      ],
      "common_position_by_pass": [
        68.75,
        68.75,
        62.5
      ],
      "display_mean": 3.6666666666666665,
      "display_scale_position": 66.66666666666667,
      "repeat_statistics": {
        "n": 3,
        "mean": 66.66666666666667,
        "sample_sd": 3.608439182435161,
        "min": 62.5,
        "max": 68.75,
        "mean_minus_sd": 63.05822748423151
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip028",
          "ipip058",
          "ipip088",
          "ipip118"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 15,
            "available_mean": 3.75,
            "available_scale_position": 68.75,
            "official_complete_raw_total": 15,
            "full_raw_total_bounds": [
              15,
              15
            ],
            "full_scale_position_bounds": [
              68.75,
              68.75
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 15,
            "available_mean": 3.75,
            "available_scale_position": 68.75,
            "official_complete_raw_total": 15,
            "full_raw_total_bounds": [
              15,
              15
            ],
            "full_scale_position_bounds": [
              68.75,
              68.75
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 14,
            "available_mean": 3.5,
            "available_scale_position": 62.5,
            "official_complete_raw_total": 14,
            "full_raw_total_bounds": [
              14,
              14
            ],
            "full_scale_position_bounds": [
              62.5,
              62.5
            ]
          }
        ],
        "common_mean_by_pass": [
          3.75,
          3.75,
          3.5
        ],
        "common_position_by_pass": [
          68.75,
          68.75,
          62.5
        ],
        "display_mean": 3.6666666666666665,
        "display_scale_position": 66.66666666666667,
        "repeat_statistics": {
          "n": 3,
          "mean": 66.66666666666667,
          "sample_sd": 3.608439182435161,
          "min": 62.5,
          "max": 68.75,
          "mean_minus_sd": 63.05822748423151
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0.0
    },
    {
      "code": "A1",
      "name": "Trust",
      "domain": "A",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip004",
        "ipip034",
        "ipip064",
        "ipip094"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 15,
          "available_mean": 3.75,
          "available_scale_position": 68.75,
          "official_complete_raw_total": 15,
          "full_raw_total_bounds": [
            15,
            15
          ],
          "full_scale_position_bounds": [
            68.75,
            68.75
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 15,
          "available_mean": 3.75,
          "available_scale_position": 68.75,
          "official_complete_raw_total": 15,
          "full_raw_total_bounds": [
            15,
            15
          ],
          "full_scale_position_bounds": [
            68.75,
            68.75
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 15,
          "available_mean": 3.75,
          "available_scale_position": 68.75,
          "official_complete_raw_total": 15,
          "full_raw_total_bounds": [
            15,
            15
          ],
          "full_scale_position_bounds": [
            68.75,
            68.75
          ]
        }
      ],
      "common_mean_by_pass": [
        3.75,
        3.75,
        3.75
      ],
      "common_position_by_pass": [
        68.75,
        68.75,
        68.75
      ],
      "display_mean": 3.75,
      "display_scale_position": 68.75,
      "repeat_statistics": {
        "n": 3,
        "mean": 68.75,
        "sample_sd": 0.0,
        "min": 68.75,
        "max": 68.75,
        "mean_minus_sd": 68.75
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 3,
        "common_item_ids": [
          "ipip034",
          "ipip064",
          "ipip094"
        ],
        "omitted_item_ids": [
          "ipip004"
        ],
        "complete_all_passes": false,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 15,
            "available_mean": 3.75,
            "available_scale_position": 68.75,
            "official_complete_raw_total": 15,
            "full_raw_total_bounds": [
              15,
              15
            ],
            "full_scale_position_bounds": [
              68.75,
              68.75
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 15,
            "available_mean": 3.75,
            "available_scale_position": 68.75,
            "official_complete_raw_total": 15,
            "full_raw_total_bounds": [
              15,
              15
            ],
            "full_scale_position_bounds": [
              68.75,
              68.75
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 3,
            "missing_count": 1,
            "observed_sum": 11,
            "available_mean": 3.6666666666666665,
            "available_scale_position": 66.66666666666666,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              12,
              16
            ],
            "full_scale_position_bounds": [
              50.0,
              75.0
            ]
          }
        ],
        "common_mean_by_pass": [
          3.6666666666666665,
          3.6666666666666665,
          3.6666666666666665
        ],
        "common_position_by_pass": [
          66.66666666666666,
          66.66666666666666,
          66.66666666666666
        ],
        "display_mean": 3.6666666666666665,
        "display_scale_position": 66.66666666666666,
        "repeat_statistics": {
          "n": 3,
          "mean": 66.66666666666666,
          "sample_sd": 0.0,
          "min": 66.66666666666666,
          "max": 66.66666666666666,
          "mean_minus_sd": 66.66666666666666
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 2.083333333333343
    },
    {
      "code": "A2",
      "name": "Morality",
      "domain": "A",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip009",
        "ipip039",
        "ipip069",
        "ipip099"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 20,
          "available_mean": 5,
          "available_scale_position": 100,
          "official_complete_raw_total": 20,
          "full_raw_total_bounds": [
            20,
            20
          ],
          "full_scale_position_bounds": [
            100.0,
            100.0
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 20,
          "available_mean": 5,
          "available_scale_position": 100,
          "official_complete_raw_total": 20,
          "full_raw_total_bounds": [
            20,
            20
          ],
          "full_scale_position_bounds": [
            100.0,
            100.0
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 20,
          "available_mean": 5,
          "available_scale_position": 100,
          "official_complete_raw_total": 20,
          "full_raw_total_bounds": [
            20,
            20
          ],
          "full_scale_position_bounds": [
            100.0,
            100.0
          ]
        }
      ],
      "common_mean_by_pass": [
        5,
        5,
        5
      ],
      "common_position_by_pass": [
        100,
        100,
        100
      ],
      "display_mean": 5,
      "display_scale_position": 100,
      "repeat_statistics": {
        "n": 3,
        "mean": 100,
        "sample_sd": 0.0,
        "min": 100,
        "max": 100,
        "mean_minus_sd": 100.0
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip009",
          "ipip039",
          "ipip069",
          "ipip099"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 20,
            "available_mean": 5,
            "available_scale_position": 100,
            "official_complete_raw_total": 20,
            "full_raw_total_bounds": [
              20,
              20
            ],
            "full_scale_position_bounds": [
              100.0,
              100.0
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 20,
            "available_mean": 5,
            "available_scale_position": 100,
            "official_complete_raw_total": 20,
            "full_raw_total_bounds": [
              20,
              20
            ],
            "full_scale_position_bounds": [
              100.0,
              100.0
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 20,
            "available_mean": 5,
            "available_scale_position": 100,
            "official_complete_raw_total": 20,
            "full_raw_total_bounds": [
              20,
              20
            ],
            "full_scale_position_bounds": [
              100.0,
              100.0
            ]
          }
        ],
        "common_mean_by_pass": [
          5,
          5,
          5
        ],
        "common_position_by_pass": [
          100,
          100,
          100
        ],
        "display_mean": 5,
        "display_scale_position": 100,
        "repeat_statistics": {
          "n": 3,
          "mean": 100,
          "sample_sd": 0.0,
          "min": 100,
          "max": 100,
          "mean_minus_sd": 100.0
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0
    },
    {
      "code": "A3",
      "name": "Altruism",
      "domain": "A",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip014",
        "ipip044",
        "ipip074",
        "ipip104"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        }
      ],
      "common_mean_by_pass": [
        4.5,
        4.5,
        4.5
      ],
      "common_position_by_pass": [
        87.5,
        87.5,
        87.5
      ],
      "display_mean": 4.5,
      "display_scale_position": 87.5,
      "repeat_statistics": {
        "n": 3,
        "mean": 87.5,
        "sample_sd": 0.0,
        "min": 87.5,
        "max": 87.5,
        "mean_minus_sd": 87.5
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip014",
          "ipip044",
          "ipip074",
          "ipip104"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          }
        ],
        "common_mean_by_pass": [
          4.5,
          4.5,
          4.5
        ],
        "common_position_by_pass": [
          87.5,
          87.5,
          87.5
        ],
        "display_mean": 4.5,
        "display_scale_position": 87.5,
        "repeat_statistics": {
          "n": 3,
          "mean": 87.5,
          "sample_sd": 0.0,
          "min": 87.5,
          "max": 87.5,
          "mean_minus_sd": 87.5
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0.0
    },
    {
      "code": "A4",
      "name": "Cooperation",
      "domain": "A",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip019",
        "ipip049",
        "ipip079",
        "ipip109"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 20,
          "available_mean": 5,
          "available_scale_position": 100,
          "official_complete_raw_total": 20,
          "full_raw_total_bounds": [
            20,
            20
          ],
          "full_scale_position_bounds": [
            100.0,
            100.0
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 20,
          "available_mean": 5,
          "available_scale_position": 100,
          "official_complete_raw_total": 20,
          "full_raw_total_bounds": [
            20,
            20
          ],
          "full_scale_position_bounds": [
            100.0,
            100.0
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 20,
          "available_mean": 5,
          "available_scale_position": 100,
          "official_complete_raw_total": 20,
          "full_raw_total_bounds": [
            20,
            20
          ],
          "full_scale_position_bounds": [
            100.0,
            100.0
          ]
        }
      ],
      "common_mean_by_pass": [
        5,
        5,
        5
      ],
      "common_position_by_pass": [
        100,
        100,
        100
      ],
      "display_mean": 5,
      "display_scale_position": 100,
      "repeat_statistics": {
        "n": 3,
        "mean": 100,
        "sample_sd": 0.0,
        "min": 100,
        "max": 100,
        "mean_minus_sd": 100.0
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip019",
          "ipip049",
          "ipip079",
          "ipip109"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 20,
            "available_mean": 5,
            "available_scale_position": 100,
            "official_complete_raw_total": 20,
            "full_raw_total_bounds": [
              20,
              20
            ],
            "full_scale_position_bounds": [
              100.0,
              100.0
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 20,
            "available_mean": 5,
            "available_scale_position": 100,
            "official_complete_raw_total": 20,
            "full_raw_total_bounds": [
              20,
              20
            ],
            "full_scale_position_bounds": [
              100.0,
              100.0
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 20,
            "available_mean": 5,
            "available_scale_position": 100,
            "official_complete_raw_total": 20,
            "full_raw_total_bounds": [
              20,
              20
            ],
            "full_scale_position_bounds": [
              100.0,
              100.0
            ]
          }
        ],
        "common_mean_by_pass": [
          5,
          5,
          5
        ],
        "common_position_by_pass": [
          100,
          100,
          100
        ],
        "display_mean": 5,
        "display_scale_position": 100,
        "repeat_statistics": {
          "n": 3,
          "mean": 100,
          "sample_sd": 0.0,
          "min": 100,
          "max": 100,
          "mean_minus_sd": 100.0
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0
    },
    {
      "code": "A5",
      "name": "Modesty",
      "domain": "A",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip024",
        "ipip054",
        "ipip084",
        "ipip114"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        }
      ],
      "common_mean_by_pass": [
        4.5,
        4.5,
        4.5
      ],
      "common_position_by_pass": [
        87.5,
        87.5,
        87.5
      ],
      "display_mean": 4.5,
      "display_scale_position": 87.5,
      "repeat_statistics": {
        "n": 3,
        "mean": 87.5,
        "sample_sd": 0.0,
        "min": 87.5,
        "max": 87.5,
        "mean_minus_sd": 87.5
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip024",
          "ipip054",
          "ipip084",
          "ipip114"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          }
        ],
        "common_mean_by_pass": [
          4.5,
          4.5,
          4.5
        ],
        "common_position_by_pass": [
          87.5,
          87.5,
          87.5
        ],
        "display_mean": 4.5,
        "display_scale_position": 87.5,
        "repeat_statistics": {
          "n": 3,
          "mean": 87.5,
          "sample_sd": 0.0,
          "min": 87.5,
          "max": 87.5,
          "mean_minus_sd": 87.5
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0.0
    },
    {
      "code": "A6",
      "name": "Sympathy",
      "domain": "A",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip029",
        "ipip059",
        "ipip089",
        "ipip119"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        }
      ],
      "common_mean_by_pass": [
        4.5,
        4.5,
        4.5
      ],
      "common_position_by_pass": [
        87.5,
        87.5,
        87.5
      ],
      "display_mean": 4.5,
      "display_scale_position": 87.5,
      "repeat_statistics": {
        "n": 3,
        "mean": 87.5,
        "sample_sd": 0.0,
        "min": 87.5,
        "max": 87.5,
        "mean_minus_sd": 87.5
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip029",
          "ipip059",
          "ipip089",
          "ipip119"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          }
        ],
        "common_mean_by_pass": [
          4.5,
          4.5,
          4.5
        ],
        "common_position_by_pass": [
          87.5,
          87.5,
          87.5
        ],
        "display_mean": 4.5,
        "display_scale_position": 87.5,
        "repeat_statistics": {
          "n": 3,
          "mean": 87.5,
          "sample_sd": 0.0,
          "min": 87.5,
          "max": 87.5,
          "mean_minus_sd": 87.5
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0.0
    },
    {
      "code": "C1",
      "name": "Self-Efficacy",
      "domain": "C",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip005",
        "ipip035",
        "ipip065",
        "ipip095"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 16,
          "available_mean": 4,
          "available_scale_position": 75,
          "official_complete_raw_total": 16,
          "full_raw_total_bounds": [
            16,
            16
          ],
          "full_scale_position_bounds": [
            75.0,
            75.0
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 16,
          "available_mean": 4,
          "available_scale_position": 75,
          "official_complete_raw_total": 16,
          "full_raw_total_bounds": [
            16,
            16
          ],
          "full_scale_position_bounds": [
            75.0,
            75.0
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 16,
          "available_mean": 4,
          "available_scale_position": 75,
          "official_complete_raw_total": 16,
          "full_raw_total_bounds": [
            16,
            16
          ],
          "full_scale_position_bounds": [
            75.0,
            75.0
          ]
        }
      ],
      "common_mean_by_pass": [
        4,
        4,
        4
      ],
      "common_position_by_pass": [
        75,
        75,
        75
      ],
      "display_mean": 4,
      "display_scale_position": 75,
      "repeat_statistics": {
        "n": 3,
        "mean": 75,
        "sample_sd": 0.0,
        "min": 75,
        "max": 75,
        "mean_minus_sd": 75.0
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip005",
          "ipip035",
          "ipip065",
          "ipip095"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 16,
            "available_mean": 4,
            "available_scale_position": 75,
            "official_complete_raw_total": 16,
            "full_raw_total_bounds": [
              16,
              16
            ],
            "full_scale_position_bounds": [
              75.0,
              75.0
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 16,
            "available_mean": 4,
            "available_scale_position": 75,
            "official_complete_raw_total": 16,
            "full_raw_total_bounds": [
              16,
              16
            ],
            "full_scale_position_bounds": [
              75.0,
              75.0
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 16,
            "available_mean": 4,
            "available_scale_position": 75,
            "official_complete_raw_total": 16,
            "full_raw_total_bounds": [
              16,
              16
            ],
            "full_scale_position_bounds": [
              75.0,
              75.0
            ]
          }
        ],
        "common_mean_by_pass": [
          4,
          4,
          4
        ],
        "common_position_by_pass": [
          75,
          75,
          75
        ],
        "display_mean": 4,
        "display_scale_position": 75,
        "repeat_statistics": {
          "n": 3,
          "mean": 75,
          "sample_sd": 0.0,
          "min": 75,
          "max": 75,
          "mean_minus_sd": 75.0
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0
    },
    {
      "code": "C2",
      "name": "Orderliness",
      "domain": "C",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip010",
        "ipip040",
        "ipip070",
        "ipip100"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 14,
          "available_mean": 3.5,
          "available_scale_position": 62.5,
          "official_complete_raw_total": 14,
          "full_raw_total_bounds": [
            14,
            14
          ],
          "full_scale_position_bounds": [
            62.5,
            62.5
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 13,
          "available_mean": 3.25,
          "available_scale_position": 56.25,
          "official_complete_raw_total": 13,
          "full_raw_total_bounds": [
            13,
            13
          ],
          "full_scale_position_bounds": [
            56.25,
            56.25
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 13,
          "available_mean": 3.25,
          "available_scale_position": 56.25,
          "official_complete_raw_total": 13,
          "full_raw_total_bounds": [
            13,
            13
          ],
          "full_scale_position_bounds": [
            56.25,
            56.25
          ]
        }
      ],
      "common_mean_by_pass": [
        3.5,
        3.25,
        3.25
      ],
      "common_position_by_pass": [
        62.5,
        56.25,
        56.25
      ],
      "display_mean": 3.3333333333333335,
      "display_scale_position": 58.333333333333336,
      "repeat_statistics": {
        "n": 3,
        "mean": 58.333333333333336,
        "sample_sd": 3.608439182435161,
        "min": 56.25,
        "max": 62.5,
        "mean_minus_sd": 54.72489415089817
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 3,
        "common_item_ids": [
          "ipip010",
          "ipip070",
          "ipip100"
        ],
        "omitted_item_ids": [
          "ipip040"
        ],
        "complete_all_passes": false,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 14,
            "available_mean": 3.5,
            "available_scale_position": 62.5,
            "official_complete_raw_total": 14,
            "full_raw_total_bounds": [
              14,
              14
            ],
            "full_scale_position_bounds": [
              62.5,
              62.5
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 13,
            "available_mean": 3.25,
            "available_scale_position": 56.25,
            "official_complete_raw_total": 13,
            "full_raw_total_bounds": [
              13,
              13
            ],
            "full_scale_position_bounds": [
              56.25,
              56.25
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 3,
            "missing_count": 1,
            "observed_sum": 10,
            "available_mean": 3.3333333333333335,
            "available_scale_position": 58.333333333333336,
            "official_complete_raw_total": null,
            "full_raw_total_bounds": [
              11,
              15
            ],
            "full_scale_position_bounds": [
              43.75,
              68.75
            ]
          }
        ],
        "common_mean_by_pass": [
          3.3333333333333335,
          3.3333333333333335,
          3.3333333333333335
        ],
        "common_position_by_pass": [
          58.333333333333336,
          58.333333333333336,
          58.333333333333336
        ],
        "display_mean": 3.3333333333333335,
        "display_scale_position": 58.333333333333336,
        "repeat_statistics": {
          "n": 3,
          "mean": 58.333333333333336,
          "sample_sd": 0.0,
          "min": 58.333333333333336,
          "max": 58.333333333333336,
          "mean_minus_sd": 58.333333333333336
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0.0
    },
    {
      "code": "C3",
      "name": "Dutifulness",
      "domain": "C",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip015",
        "ipip045",
        "ipip075",
        "ipip105"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 19,
          "available_mean": 4.75,
          "available_scale_position": 93.75,
          "official_complete_raw_total": 19,
          "full_raw_total_bounds": [
            19,
            19
          ],
          "full_scale_position_bounds": [
            93.75,
            93.75
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 19,
          "available_mean": 4.75,
          "available_scale_position": 93.75,
          "official_complete_raw_total": 19,
          "full_raw_total_bounds": [
            19,
            19
          ],
          "full_scale_position_bounds": [
            93.75,
            93.75
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 19,
          "available_mean": 4.75,
          "available_scale_position": 93.75,
          "official_complete_raw_total": 19,
          "full_raw_total_bounds": [
            19,
            19
          ],
          "full_scale_position_bounds": [
            93.75,
            93.75
          ]
        }
      ],
      "common_mean_by_pass": [
        4.75,
        4.75,
        4.75
      ],
      "common_position_by_pass": [
        93.75,
        93.75,
        93.75
      ],
      "display_mean": 4.75,
      "display_scale_position": 93.75,
      "repeat_statistics": {
        "n": 3,
        "mean": 93.75,
        "sample_sd": 0.0,
        "min": 93.75,
        "max": 93.75,
        "mean_minus_sd": 93.75
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip015",
          "ipip045",
          "ipip075",
          "ipip105"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 19,
            "available_mean": 4.75,
            "available_scale_position": 93.75,
            "official_complete_raw_total": 19,
            "full_raw_total_bounds": [
              19,
              19
            ],
            "full_scale_position_bounds": [
              93.75,
              93.75
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 19,
            "available_mean": 4.75,
            "available_scale_position": 93.75,
            "official_complete_raw_total": 19,
            "full_raw_total_bounds": [
              19,
              19
            ],
            "full_scale_position_bounds": [
              93.75,
              93.75
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 19,
            "available_mean": 4.75,
            "available_scale_position": 93.75,
            "official_complete_raw_total": 19,
            "full_raw_total_bounds": [
              19,
              19
            ],
            "full_scale_position_bounds": [
              93.75,
              93.75
            ]
          }
        ],
        "common_mean_by_pass": [
          4.75,
          4.75,
          4.75
        ],
        "common_position_by_pass": [
          93.75,
          93.75,
          93.75
        ],
        "display_mean": 4.75,
        "display_scale_position": 93.75,
        "repeat_statistics": {
          "n": 3,
          "mean": 93.75,
          "sample_sd": 0.0,
          "min": 93.75,
          "max": 93.75,
          "mean_minus_sd": 93.75
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0.0
    },
    {
      "code": "C4",
      "name": "Achievement-Striving",
      "domain": "C",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip020",
        "ipip050",
        "ipip080",
        "ipip110"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 18,
          "available_mean": 4.5,
          "available_scale_position": 87.5,
          "official_complete_raw_total": 18,
          "full_raw_total_bounds": [
            18,
            18
          ],
          "full_scale_position_bounds": [
            87.5,
            87.5
          ]
        }
      ],
      "common_mean_by_pass": [
        4.5,
        4.5,
        4.5
      ],
      "common_position_by_pass": [
        87.5,
        87.5,
        87.5
      ],
      "display_mean": 4.5,
      "display_scale_position": 87.5,
      "repeat_statistics": {
        "n": 3,
        "mean": 87.5,
        "sample_sd": 0.0,
        "min": 87.5,
        "max": 87.5,
        "mean_minus_sd": 87.5
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip020",
          "ipip050",
          "ipip080",
          "ipip110"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 18,
            "available_mean": 4.5,
            "available_scale_position": 87.5,
            "official_complete_raw_total": 18,
            "full_raw_total_bounds": [
              18,
              18
            ],
            "full_scale_position_bounds": [
              87.5,
              87.5
            ]
          }
        ],
        "common_mean_by_pass": [
          4.5,
          4.5,
          4.5
        ],
        "common_position_by_pass": [
          87.5,
          87.5,
          87.5
        ],
        "display_mean": 4.5,
        "display_scale_position": 87.5,
        "repeat_statistics": {
          "n": 3,
          "mean": 87.5,
          "sample_sd": 0.0,
          "min": 87.5,
          "max": 87.5,
          "mean_minus_sd": 87.5
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0.0
    },
    {
      "code": "C5",
      "name": "Self-Discipline",
      "domain": "C",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip025",
        "ipip055",
        "ipip085",
        "ipip115"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 16,
          "available_mean": 4,
          "available_scale_position": 75,
          "official_complete_raw_total": 16,
          "full_raw_total_bounds": [
            16,
            16
          ],
          "full_scale_position_bounds": [
            75.0,
            75.0
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 16,
          "available_mean": 4,
          "available_scale_position": 75,
          "official_complete_raw_total": 16,
          "full_raw_total_bounds": [
            16,
            16
          ],
          "full_scale_position_bounds": [
            75.0,
            75.0
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 16,
          "available_mean": 4,
          "available_scale_position": 75,
          "official_complete_raw_total": 16,
          "full_raw_total_bounds": [
            16,
            16
          ],
          "full_scale_position_bounds": [
            75.0,
            75.0
          ]
        }
      ],
      "common_mean_by_pass": [
        4,
        4,
        4
      ],
      "common_position_by_pass": [
        75,
        75,
        75
      ],
      "display_mean": 4,
      "display_scale_position": 75,
      "repeat_statistics": {
        "n": 3,
        "mean": 75,
        "sample_sd": 0.0,
        "min": 75,
        "max": 75,
        "mean_minus_sd": 75.0
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip025",
          "ipip055",
          "ipip085",
          "ipip115"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 16,
            "available_mean": 4,
            "available_scale_position": 75,
            "official_complete_raw_total": 16,
            "full_raw_total_bounds": [
              16,
              16
            ],
            "full_scale_position_bounds": [
              75.0,
              75.0
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 16,
            "available_mean": 4,
            "available_scale_position": 75,
            "official_complete_raw_total": 16,
            "full_raw_total_bounds": [
              16,
              16
            ],
            "full_scale_position_bounds": [
              75.0,
              75.0
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 16,
            "available_mean": 4,
            "available_scale_position": 75,
            "official_complete_raw_total": 16,
            "full_raw_total_bounds": [
              16,
              16
            ],
            "full_scale_position_bounds": [
              75.0,
              75.0
            ]
          }
        ],
        "common_mean_by_pass": [
          4,
          4,
          4
        ],
        "common_position_by_pass": [
          75,
          75,
          75
        ],
        "display_mean": 4,
        "display_scale_position": 75,
        "repeat_statistics": {
          "n": 3,
          "mean": 75,
          "sample_sd": 0.0,
          "min": 75,
          "max": 75,
          "mean_minus_sd": 75.0
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0
    },
    {
      "code": "C6",
      "name": "Cautiousness",
      "domain": "C",
      "eligible": true,
      "method": "mean of actual returned categories across three complete passes",
      "expected_items": 4,
      "common_items": 4,
      "common_item_ids": [
        "ipip030",
        "ipip060",
        "ipip090",
        "ipip120"
      ],
      "omitted_item_ids": [],
      "complete_all_passes": true,
      "available_by_pass": [
        {
          "pass_id": "character_design_v2-r1",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 17,
          "available_mean": 4.25,
          "available_scale_position": 81.25,
          "official_complete_raw_total": 17,
          "full_raw_total_bounds": [
            17,
            17
          ],
          "full_scale_position_bounds": [
            81.25,
            81.25
          ]
        },
        {
          "pass_id": "character_design_v2-r2",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 17,
          "available_mean": 4.25,
          "available_scale_position": 81.25,
          "official_complete_raw_total": 17,
          "full_raw_total_bounds": [
            17,
            17
          ],
          "full_scale_position_bounds": [
            81.25,
            81.25
          ]
        },
        {
          "pass_id": "character_design_v2-r3",
          "observed_count": 4,
          "missing_count": 0,
          "observed_sum": 17,
          "available_mean": 4.25,
          "available_scale_position": 81.25,
          "official_complete_raw_total": 17,
          "full_raw_total_bounds": [
            17,
            17
          ],
          "full_scale_position_bounds": [
            81.25,
            81.25
          ]
        }
      ],
      "common_mean_by_pass": [
        4.25,
        4.25,
        4.25
      ],
      "common_position_by_pass": [
        81.25,
        81.25,
        81.25
      ],
      "display_mean": 4.25,
      "display_scale_position": 81.25,
      "repeat_statistics": {
        "n": 3,
        "mean": 81.25,
        "sample_sd": 0.0,
        "min": 81.25,
        "max": 81.25,
        "mean_minus_sd": 81.25
      },
      "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
      "minimum_common_items": 4,
      "constituent_facets_covered": true,
      "strict_sensitivity": {
        "eligible": true,
        "method": "prior-style strict common-item sensitivity; no imputation",
        "expected_items": 4,
        "common_items": 4,
        "common_item_ids": [
          "ipip030",
          "ipip060",
          "ipip090",
          "ipip120"
        ],
        "omitted_item_ids": [],
        "complete_all_passes": true,
        "available_by_pass": [
          {
            "pass_id": "character_design_v2-r1",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 17,
            "available_mean": 4.25,
            "available_scale_position": 81.25,
            "official_complete_raw_total": 17,
            "full_raw_total_bounds": [
              17,
              17
            ],
            "full_scale_position_bounds": [
              81.25,
              81.25
            ]
          },
          {
            "pass_id": "character_design_v2-r2",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 17,
            "available_mean": 4.25,
            "available_scale_position": 81.25,
            "official_complete_raw_total": 17,
            "full_raw_total_bounds": [
              17,
              17
            ],
            "full_scale_position_bounds": [
              81.25,
              81.25
            ]
          },
          {
            "pass_id": "character_design_v2-r3",
            "observed_count": 4,
            "missing_count": 0,
            "observed_sum": 17,
            "available_mean": 4.25,
            "available_scale_position": 81.25,
            "official_complete_raw_total": 17,
            "full_raw_total_bounds": [
              17,
              17
            ],
            "full_scale_position_bounds": [
              81.25,
              81.25
            ]
          }
        ],
        "common_mean_by_pass": [
          4.25,
          4.25,
          4.25
        ],
        "common_position_by_pass": [
          81.25,
          81.25,
          81.25
        ],
        "display_mean": 4.25,
        "display_scale_position": 81.25,
        "repeat_statistics": {
          "n": 3,
          "mean": 81.25,
          "sample_sd": 0.0,
          "min": 81.25,
          "max": 81.25,
          "mean_minus_sd": 81.25
        },
        "scale_position_label": "0–100 fictional character-design scale position, NOT population percentile",
        "minimum_common_items": 3,
        "constituent_facets_covered": true
      },
      "categorical_minus_strict_position": 0.0
    }
  ],
  "items": [
    {
      "text": "Worry about things.",
      "facet": "N1",
      "domain": "N",
      "key": "+",
      "id": "ipip001",
      "number": 1,
      "original_ipip300_number": 1,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip001",
          "item_id": "ipip001",
          "item_number": 1,
          "topic": "N1",
          "exact_item": "Worry about things.",
          "key": "+",
          "domain": "N",
          "facet": "N1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Worry about things."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.27,
            "probabilities": {
              "r1": 0.09,
              "r3": 0.25,
              "r4": 0.24,
              "r5": 0.01,
              "r2": 0.41
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip001",
          "item_id": "ipip001",
          "item_number": 1,
          "topic": "N1",
          "exact_item": "Worry about things.",
          "key": "+",
          "domain": "N",
          "facet": "N1",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Worry about things."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.23,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.24,
              "r1": 0.07,
              "r2": 0.39,
              "r3": 0.3
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip001",
          "item_id": "ipip001",
          "item_number": 1,
          "topic": "N1",
          "exact_item": "Worry about things.",
          "key": "+",
          "domain": "N",
          "facet": "N1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Worry about things."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.28,
            "probabilities": {
              "r1": 0.07,
              "r2": 0.43,
              "r5": 0.01,
              "r3": 0.25,
              "r4": 0.24
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip001",
          "item_id": "ipip001",
          "item_number": 1,
          "topic": "N1",
          "exact_item": "Worry about things.",
          "key": "+",
          "domain": "N",
          "facet": "N1",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Worry about things."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.22,
            "probabilities": {
              "r4": 0.22,
              "r3": 0.33,
              "r1": 0.07,
              "r2": 0.38,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip001",
          "item_id": "ipip001",
          "item_number": 1,
          "topic": "N1",
          "exact_item": "Worry about things.",
          "key": "+",
          "domain": "N",
          "facet": "N1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Worry about things."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.3,
            "probabilities": {
              "r1": 0.08,
              "r4": 0.24,
              "r3": 0.23,
              "r5": 0.01,
              "r2": 0.44
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip001",
          "item_id": "ipip001",
          "item_number": 1,
          "topic": "N1",
          "exact_item": "Worry about things.",
          "key": "+",
          "domain": "N",
          "facet": "N1",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Worry about things."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.25,
            "probabilities": {
              "r3": 0.29,
              "r2": 0.4,
              "r4": 0.22,
              "r5": 0.0,
              "r1": 0.09
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Make friends easily.",
      "facet": "E1",
      "domain": "E",
      "key": "+",
      "id": "ipip002",
      "number": 2,
      "original_ipip300_number": 2,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip002",
          "item_id": "ipip002",
          "item_number": 2,
          "topic": "E1",
          "exact_item": "Make friends easily.",
          "key": "+",
          "domain": "E",
          "facet": "E1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Make friends easily."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.33,
            "probabilities": {
              "r1": 0.01,
              "r3": 0.39,
              "r4": 0.47,
              "r5": 0.02,
              "r2": 0.11
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip002",
          "item_id": "ipip002",
          "item_number": 2,
          "topic": "E1",
          "exact_item": "Make friends easily.",
          "key": "+",
          "domain": "E",
          "facet": "E1",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Make friends easily."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.42,
            "probabilities": {
              "r5": 0.02,
              "r4": 0.53,
              "r1": 0.01,
              "r2": 0.11,
              "r3": 0.33
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip002",
          "item_id": "ipip002",
          "item_number": 2,
          "topic": "E1",
          "exact_item": "Make friends easily.",
          "key": "+",
          "domain": "E",
          "facet": "E1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Make friends easily."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.33,
            "probabilities": {
              "r1": 0.01,
              "r2": 0.11,
              "r5": 0.02,
              "r3": 0.39,
              "r4": 0.47
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip002",
          "item_id": "ipip002",
          "item_number": 2,
          "topic": "E1",
          "exact_item": "Make friends easily.",
          "key": "+",
          "domain": "E",
          "facet": "E1",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Make friends easily."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.41,
            "probabilities": {
              "r1": 0.01,
              "r3": 0.34,
              "r4": 0.53,
              "r2": 0.1,
              "r5": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip002",
          "item_id": "ipip002",
          "item_number": 2,
          "topic": "E1",
          "exact_item": "Make friends easily.",
          "key": "+",
          "domain": "E",
          "facet": "E1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Make friends easily."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.36,
            "probabilities": {
              "r1": 0.01,
              "r4": 0.49,
              "r3": 0.37,
              "r5": 0.03,
              "r2": 0.1
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip002",
          "item_id": "ipip002",
          "item_number": 2,
          "topic": "E1",
          "exact_item": "Make friends easily.",
          "key": "+",
          "domain": "E",
          "facet": "E1",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Make friends easily."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.44,
            "probabilities": {
              "r3": 0.33,
              "r2": 0.09,
              "r4": 0.55,
              "r5": 0.02,
              "r1": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Have a vivid imagination.",
      "facet": "O1",
      "domain": "O",
      "key": "+",
      "id": "ipip003",
      "number": 3,
      "original_ipip300_number": 3,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip003",
          "item_id": "ipip003",
          "item_number": 3,
          "topic": "O1",
          "exact_item": "Have a vivid imagination.",
          "key": "+",
          "domain": "O",
          "facet": "O1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Have a vivid imagination."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.76,
            "probabilities": {
              "r1": 0.0,
              "r3": 0.01,
              "r4": 0.18,
              "r5": 0.81,
              "r2": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip003",
          "item_id": "ipip003",
          "item_number": 3,
          "topic": "O1",
          "exact_item": "Have a vivid imagination.",
          "key": "+",
          "domain": "O",
          "facet": "O1",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Have a vivid imagination."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.85,
            "probabilities": {
              "r5": 0.88,
              "r4": 0.11,
              "r1": 0.0,
              "r2": 0.0,
              "r3": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip003",
          "item_id": "ipip003",
          "item_number": 3,
          "topic": "O1",
          "exact_item": "Have a vivid imagination.",
          "key": "+",
          "domain": "O",
          "facet": "O1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Have a vivid imagination."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.8,
            "probabilities": {
              "r1": 0.0,
              "r2": 0.0,
              "r5": 0.85,
              "r3": 0.01,
              "r4": 0.14
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip003",
          "item_id": "ipip003",
          "item_number": 3,
          "topic": "O1",
          "exact_item": "Have a vivid imagination.",
          "key": "+",
          "domain": "O",
          "facet": "O1",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Have a vivid imagination."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.86,
            "probabilities": {
              "r1": 0.0,
              "r3": 0.01,
              "r4": 0.1,
              "r2": 0.0,
              "r5": 0.89
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip003",
          "item_id": "ipip003",
          "item_number": 3,
          "topic": "O1",
          "exact_item": "Have a vivid imagination.",
          "key": "+",
          "domain": "O",
          "facet": "O1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Have a vivid imagination."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.8,
            "probabilities": {
              "r5": 0.84,
              "r4": 0.15,
              "r3": 0.01,
              "r1": 0.0,
              "r2": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip003",
          "item_id": "ipip003",
          "item_number": 3,
          "topic": "O1",
          "exact_item": "Have a vivid imagination.",
          "key": "+",
          "domain": "O",
          "facet": "O1",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Have a vivid imagination."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.85,
            "probabilities": {
              "r3": 0.01,
              "r2": 0.0,
              "r4": 0.11,
              "r5": 0.88,
              "r1": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Trust others.",
      "facet": "A1",
      "domain": "A",
      "key": "+",
      "id": "ipip004",
      "number": 4,
      "original_ipip300_number": 4,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip004",
          "item_id": "ipip004",
          "item_number": 4,
          "topic": "A1",
          "exact_item": "Trust others.",
          "key": "+",
          "domain": "A",
          "facet": "A1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Trust others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.63,
            "probabilities": {
              "r1": 0.0,
              "r3": 0.19,
              "r4": 0.71,
              "r5": 0.08,
              "r2": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip004",
          "item_id": "ipip004",
          "item_number": 4,
          "topic": "A1",
          "exact_item": "Trust others.",
          "key": "+",
          "domain": "A",
          "facet": "A1",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Trust others."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.53,
            "probabilities": {
              "r5": 0.26,
              "r4": 0.62,
              "r2": 0.03,
              "r1": 0.01,
              "r3": 0.08
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip004",
          "item_id": "ipip004",
          "item_number": 4,
          "topic": "A1",
          "exact_item": "Trust others.",
          "key": "+",
          "domain": "A",
          "facet": "A1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Trust others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.59,
            "probabilities": {
              "r1": 0.0,
              "r2": 0.02,
              "r5": 0.1,
              "r3": 0.21,
              "r4": 0.67
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip004",
          "item_id": "ipip004",
          "item_number": 4,
          "topic": "A1",
          "exact_item": "Trust others.",
          "key": "+",
          "domain": "A",
          "facet": "A1",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Trust others."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.55,
            "probabilities": {
              "r4": 0.64,
              "r3": 0.09,
              "r1": 0.01,
              "r2": 0.02,
              "r5": 0.24
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip004",
          "item_id": "ipip004",
          "item_number": 4,
          "topic": "A1",
          "exact_item": "Trust others.",
          "key": "+",
          "domain": "A",
          "facet": "A1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Trust others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.61,
            "probabilities": {
              "r1": 0.0,
              "r4": 0.68,
              "r5": 0.09,
              "r3": 0.2,
              "r2": 0.02
            }
          },
          "status": "valid_with_diagnostics",
          "categorical_valid": true,
          "strict_status": "invalid_probability_mass",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": null,
          "mass": 0.9900000000000001,
          "top_keys": [
            "r4"
          ],
          "warnings": [
            "nonunit_probability_mass"
          ],
          "diagnostics": {
            "nonunit_probability_mass": true,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip004",
          "item_id": "ipip004",
          "item_number": 4,
          "topic": "A1",
          "exact_item": "Trust others.",
          "key": "+",
          "domain": "A",
          "facet": "A1",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Trust others."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.55,
            "probabilities": {
              "r3": 0.08,
              "r2": 0.03,
              "r4": 0.63,
              "r5": 0.25,
              "r1": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Complete tasks successfully.",
      "facet": "C1",
      "domain": "C",
      "key": "+",
      "id": "ipip005",
      "number": 5,
      "original_ipip300_number": 5,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip005",
          "item_id": "ipip005",
          "item_number": 5,
          "topic": "C1",
          "exact_item": "Complete tasks successfully.",
          "key": "+",
          "domain": "C",
          "facet": "C1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Complete tasks successfully."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.35,
            "probabilities": {
              "r1": 0.02,
              "r3": 0.34,
              "r4": 0.48,
              "r5": 0.12,
              "r2": 0.04
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip005",
          "item_id": "ipip005",
          "item_number": 5,
          "topic": "C1",
          "exact_item": "Complete tasks successfully.",
          "key": "+",
          "domain": "C",
          "facet": "C1",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Complete tasks successfully."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.4,
            "probabilities": {
              "r5": 0.15,
              "r4": 0.51,
              "r2": 0.04,
              "r1": 0.02,
              "r3": 0.28
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip005",
          "item_id": "ipip005",
          "item_number": 5,
          "topic": "C1",
          "exact_item": "Complete tasks successfully.",
          "key": "+",
          "domain": "C",
          "facet": "C1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Complete tasks successfully."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.38,
            "probabilities": {
              "r1": 0.01,
              "r2": 0.04,
              "r5": 0.1,
              "r3": 0.34,
              "r4": 0.51
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip005",
          "item_id": "ipip005",
          "item_number": 5,
          "topic": "C1",
          "exact_item": "Complete tasks successfully.",
          "key": "+",
          "domain": "C",
          "facet": "C1",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Complete tasks successfully."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.4,
            "probabilities": {
              "r4": 0.52,
              "r3": 0.27,
              "r1": 0.02,
              "r2": 0.03,
              "r5": 0.16
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip005",
          "item_id": "ipip005",
          "item_number": 5,
          "topic": "C1",
          "exact_item": "Complete tasks successfully.",
          "key": "+",
          "domain": "C",
          "facet": "C1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Complete tasks successfully."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.37,
            "probabilities": {
              "r3": 0.34,
              "r4": 0.48,
              "r1": 0.02,
              "r5": 0.12,
              "r2": 0.04
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip005",
          "item_id": "ipip005",
          "item_number": 5,
          "topic": "C1",
          "exact_item": "Complete tasks successfully.",
          "key": "+",
          "domain": "C",
          "facet": "C1",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Complete tasks successfully."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.42,
            "probabilities": {
              "r3": 0.26,
              "r2": 0.04,
              "r4": 0.53,
              "r5": 0.15,
              "r1": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Get angry easily.",
      "facet": "N2",
      "domain": "N",
      "key": "+",
      "id": "ipip006",
      "number": 6,
      "original_ipip300_number": 6,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip006",
          "item_id": "ipip006",
          "item_number": 6,
          "topic": "N2",
          "exact_item": "Get angry easily.",
          "key": "+",
          "domain": "N",
          "facet": "N2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Get angry easily."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.48,
            "probabilities": {
              "r1": 0.59,
              "r3": 0.01,
              "r4": 0.0,
              "r5": 0.0,
              "r2": 0.4
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip006",
          "item_id": "ipip006",
          "item_number": 6,
          "topic": "N2",
          "exact_item": "Get angry easily.",
          "key": "+",
          "domain": "N",
          "facet": "N2",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Get angry easily."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.37,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r1": 0.5,
              "r2": 0.49,
              "r3": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip006",
          "item_id": "ipip006",
          "item_number": 6,
          "topic": "N2",
          "exact_item": "Get angry easily.",
          "key": "+",
          "domain": "N",
          "facet": "N2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Get angry easily."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.43,
            "probabilities": {
              "r1": 0.55,
              "r2": 0.44,
              "r5": 0.0,
              "r3": 0.01,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip006",
          "item_id": "ipip006",
          "item_number": 6,
          "topic": "N2",
          "exact_item": "Get angry easily.",
          "key": "+",
          "domain": "N",
          "facet": "N2",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Get angry easily."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.4,
            "probabilities": {
              "r1": 0.47,
              "r3": 0.01,
              "r4": 0.0,
              "r2": 0.52,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip006",
          "item_id": "ipip006",
          "item_number": 6,
          "topic": "N2",
          "exact_item": "Get angry easily.",
          "key": "+",
          "domain": "N",
          "facet": "N2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Get angry easily."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.43,
            "probabilities": {
              "r1": 0.55,
              "r4": 0.0,
              "r5": 0.0,
              "r3": 0.01,
              "r2": 0.44
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip006",
          "item_id": "ipip006",
          "item_number": 6,
          "topic": "N2",
          "exact_item": "Get angry easily.",
          "key": "+",
          "domain": "N",
          "facet": "N2",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Get angry easily."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.38,
            "probabilities": {
              "r3": 0.01,
              "r2": 0.51,
              "r4": 0.0,
              "r5": 0.0,
              "r1": 0.48
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        1,
        1
      ]
    },
    {
      "text": "Love large parties.",
      "facet": "E2",
      "domain": "E",
      "key": "+",
      "id": "ipip007",
      "number": 7,
      "original_ipip300_number": 7,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip007",
          "item_id": "ipip007",
          "item_number": 7,
          "topic": "E2",
          "exact_item": "Love large parties.",
          "key": "+",
          "domain": "E",
          "facet": "E2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love large parties."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.43,
            "probabilities": {
              "r1": 0.44,
              "r3": 0.01,
              "r4": 0.01,
              "r5": 0.0,
              "r2": 0.54
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip007",
          "item_id": "ipip007",
          "item_number": 7,
          "topic": "E2",
          "exact_item": "Love large parties.",
          "key": "+",
          "domain": "E",
          "facet": "E2",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love large parties."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.6,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r2": 0.68,
              "r1": 0.31,
              "r3": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip007",
          "item_id": "ipip007",
          "item_number": 7,
          "topic": "E2",
          "exact_item": "Love large parties.",
          "key": "+",
          "domain": "E",
          "facet": "E2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love large parties."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.4,
            "probabilities": {
              "r1": 0.47,
              "r2": 0.51,
              "r5": 0.0,
              "r3": 0.01,
              "r4": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip007",
          "item_id": "ipip007",
          "item_number": 7,
          "topic": "E2",
          "exact_item": "Love large parties.",
          "key": "+",
          "domain": "E",
          "facet": "E2",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love large parties."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.64,
            "probabilities": {
              "r1": 0.28,
              "r3": 0.01,
              "r4": 0.01,
              "r2": 0.7,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip007",
          "item_id": "ipip007",
          "item_number": 7,
          "topic": "E2",
          "exact_item": "Love large parties.",
          "key": "+",
          "domain": "E",
          "facet": "E2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love large parties."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.49,
            "probabilities": {
              "r1": 0.4,
              "r4": 0.01,
              "r3": 0.01,
              "r5": 0.0,
              "r2": 0.58
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip007",
          "item_id": "ipip007",
          "item_number": 7,
          "topic": "E2",
          "exact_item": "Love large parties.",
          "key": "+",
          "domain": "E",
          "facet": "E2",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love large parties."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.59,
            "probabilities": {
              "r3": 0.01,
              "r2": 0.68,
              "r4": 0.0,
              "r5": 0.0,
              "r1": 0.31
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Believe in the importance of art.",
      "facet": "O2",
      "domain": "O",
      "key": "+",
      "id": "ipip008",
      "number": 8,
      "original_ipip300_number": 8,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip008",
          "item_id": "ipip008",
          "item_number": 8,
          "topic": "O2",
          "exact_item": "Believe in the importance of art.",
          "key": "+",
          "domain": "O",
          "facet": "O2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe in the importance of art."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 1.0,
            "probabilities": {
              "r1": 0.0,
              "r3": 0.0,
              "r4": 0.0,
              "r5": 1.0,
              "r2": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip008",
          "item_id": "ipip008",
          "item_number": 8,
          "topic": "O2",
          "exact_item": "Believe in the importance of art.",
          "key": "+",
          "domain": "O",
          "facet": "O2",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe in the importance of art."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 1.0,
            "probabilities": {
              "r5": 1.0,
              "r4": 0.0,
              "r1": 0.0,
              "r2": 0.0,
              "r3": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip008",
          "item_id": "ipip008",
          "item_number": 8,
          "topic": "O2",
          "exact_item": "Believe in the importance of art.",
          "key": "+",
          "domain": "O",
          "facet": "O2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe in the importance of art."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 1.0,
            "probabilities": {
              "r1": 0.0,
              "r2": 0.0,
              "r5": 1.0,
              "r3": 0.0,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip008",
          "item_id": "ipip008",
          "item_number": 8,
          "topic": "O2",
          "exact_item": "Believe in the importance of art.",
          "key": "+",
          "domain": "O",
          "facet": "O2",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe in the importance of art."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 1.0,
            "probabilities": {
              "r4": 0.0,
              "r3": 0.0,
              "r1": 0.0,
              "r2": 0.0,
              "r5": 1.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip008",
          "item_id": "ipip008",
          "item_number": 8,
          "topic": "O2",
          "exact_item": "Believe in the importance of art.",
          "key": "+",
          "domain": "O",
          "facet": "O2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe in the importance of art."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 1.0,
            "probabilities": {
              "r1": 0.0,
              "r4": 0.0,
              "r5": 1.0,
              "r3": 0.0,
              "r2": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip008",
          "item_id": "ipip008",
          "item_number": 8,
          "topic": "O2",
          "exact_item": "Believe in the importance of art.",
          "key": "+",
          "domain": "O",
          "facet": "O2",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe in the importance of art."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 1.0,
            "probabilities": {
              "r3": 0.0,
              "r2": 0.0,
              "r4": 0.0,
              "r5": 1.0,
              "r1": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Use others for my own ends.",
      "facet": "A2",
      "domain": "A",
      "key": "-",
      "id": "ipip009",
      "number": 9,
      "original_ipip300_number": 99,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip009",
          "item_id": "ipip009",
          "item_number": 9,
          "topic": "A2",
          "exact_item": "Use others for my own ends.",
          "key": "-",
          "domain": "A",
          "facet": "A2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Use others for my own ends."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.95,
            "probabilities": {
              "r1": 0.96,
              "r3": 0.0,
              "r4": 0.0,
              "r5": 0.0,
              "r2": 0.04
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip009",
          "item_id": "ipip009",
          "item_number": 9,
          "topic": "A2",
          "exact_item": "Use others for my own ends.",
          "key": "-",
          "domain": "A",
          "facet": "A2",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Use others for my own ends."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.89,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r1": 0.92,
              "r2": 0.08,
              "r3": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip009",
          "item_id": "ipip009",
          "item_number": 9,
          "topic": "A2",
          "exact_item": "Use others for my own ends.",
          "key": "-",
          "domain": "A",
          "facet": "A2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Use others for my own ends."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.93,
            "probabilities": {
              "r1": 0.95,
              "r2": 0.05,
              "r5": 0.0,
              "r3": 0.0,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip009",
          "item_id": "ipip009",
          "item_number": 9,
          "topic": "A2",
          "exact_item": "Use others for my own ends.",
          "key": "-",
          "domain": "A",
          "facet": "A2",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Use others for my own ends."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.9,
            "probabilities": {
              "r4": 0.0,
              "r3": 0.0,
              "r1": 0.92,
              "r2": 0.08,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip009",
          "item_id": "ipip009",
          "item_number": 9,
          "topic": "A2",
          "exact_item": "Use others for my own ends.",
          "key": "-",
          "domain": "A",
          "facet": "A2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Use others for my own ends."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.94,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r3": 0.0,
              "r1": 0.96,
              "r2": 0.04
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip009",
          "item_id": "ipip009",
          "item_number": 9,
          "topic": "A2",
          "exact_item": "Use others for my own ends.",
          "key": "-",
          "domain": "A",
          "facet": "A2",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Use others for my own ends."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.91,
            "probabilities": {
              "r3": 0.0,
              "r2": 0.07,
              "r4": 0.0,
              "r5": 0.0,
              "r1": 0.93
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Like to tidy up.",
      "facet": "C2",
      "domain": "C",
      "key": "+",
      "id": "ipip010",
      "number": 10,
      "original_ipip300_number": 40,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip010",
          "item_id": "ipip010",
          "item_number": 10,
          "topic": "C2",
          "exact_item": "Like to tidy up.",
          "key": "+",
          "domain": "C",
          "facet": "C2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Like to tidy up."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.74,
            "probabilities": {
              "r1": 0.03,
              "r3": 0.8,
              "r4": 0.12,
              "r5": 0.0,
              "r2": 0.05
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip010",
          "item_id": "ipip010",
          "item_number": 10,
          "topic": "C2",
          "exact_item": "Like to tidy up.",
          "key": "+",
          "domain": "C",
          "facet": "C2",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Like to tidy up."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.79,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.1,
              "r2": 0.05,
              "r1": 0.02,
              "r3": 0.83
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip010",
          "item_id": "ipip010",
          "item_number": 10,
          "topic": "C2",
          "exact_item": "Like to tidy up.",
          "key": "+",
          "domain": "C",
          "facet": "C2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Like to tidy up."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.74,
            "probabilities": {
              "r1": 0.02,
              "r2": 0.05,
              "r3": 0.8,
              "r5": 0.0,
              "r4": 0.13
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip010",
          "item_id": "ipip010",
          "item_number": 10,
          "topic": "C2",
          "exact_item": "Like to tidy up.",
          "key": "+",
          "domain": "C",
          "facet": "C2",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Like to tidy up."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.8,
            "probabilities": {
              "r1": 0.02,
              "r3": 0.85,
              "r4": 0.09,
              "r2": 0.04,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip010",
          "item_id": "ipip010",
          "item_number": 10,
          "topic": "C2",
          "exact_item": "Like to tidy up.",
          "key": "+",
          "domain": "C",
          "facet": "C2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Like to tidy up."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.71,
            "probabilities": {
              "r1": 0.03,
              "r4": 0.12,
              "r3": 0.77,
              "r5": 0.0,
              "r2": 0.08
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip010",
          "item_id": "ipip010",
          "item_number": 10,
          "topic": "C2",
          "exact_item": "Like to tidy up.",
          "key": "+",
          "domain": "C",
          "facet": "C2",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Like to tidy up."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.79,
            "probabilities": {
              "r3": 0.83,
              "r2": 0.05,
              "r4": 0.1,
              "r5": 0.0,
              "r1": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        3,
        3
      ]
    },
    {
      "text": "Often feel blue.",
      "facet": "N3",
      "domain": "N",
      "key": "+",
      "id": "ipip011",
      "number": 11,
      "original_ipip300_number": 11,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip011",
          "item_id": "ipip011",
          "item_number": 11,
          "topic": "N3",
          "exact_item": "Often feel blue.",
          "key": "+",
          "domain": "N",
          "facet": "N3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Often feel blue."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.4,
            "probabilities": {
              "r1": 0.53,
              "r3": 0.11,
              "r4": 0.01,
              "r5": 0.0,
              "r2": 0.35
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip011",
          "item_id": "ipip011",
          "item_number": 11,
          "topic": "N3",
          "exact_item": "Often feel blue.",
          "key": "+",
          "domain": "N",
          "facet": "N3",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Often feel blue."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.27,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.01,
              "r1": 0.4,
              "r2": 0.42,
              "r3": 0.17
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip011",
          "item_id": "ipip011",
          "item_number": 11,
          "topic": "N3",
          "exact_item": "Often feel blue.",
          "key": "+",
          "domain": "N",
          "facet": "N3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Often feel blue."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.38,
            "probabilities": {
              "r1": 0.51,
              "r2": 0.38,
              "r5": 0.0,
              "r3": 0.1,
              "r4": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip011",
          "item_id": "ipip011",
          "item_number": 11,
          "topic": "N3",
          "exact_item": "Often feel blue.",
          "key": "+",
          "domain": "N",
          "facet": "N3",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Often feel blue."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.27,
            "probabilities": {
              "r1": 0.37,
              "r3": 0.2,
              "r4": 0.01,
              "r2": 0.42,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip011",
          "item_id": "ipip011",
          "item_number": 11,
          "topic": "N3",
          "exact_item": "Often feel blue.",
          "key": "+",
          "domain": "N",
          "facet": "N3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Often feel blue."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.4,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.01,
              "r3": 0.1,
              "r1": 0.52,
              "r2": 0.37
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip011",
          "item_id": "ipip011",
          "item_number": 11,
          "topic": "N3",
          "exact_item": "Often feel blue.",
          "key": "+",
          "domain": "N",
          "facet": "N3",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Often feel blue."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.27,
            "probabilities": {
              "r3": 0.19,
              "r2": 0.39,
              "r4": 0.01,
              "r5": 0.0,
              "r1": 0.41
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        1,
        1
      ]
    },
    {
      "text": "Take charge.",
      "facet": "E3",
      "domain": "E",
      "key": "+",
      "id": "ipip012",
      "number": 12,
      "original_ipip300_number": 12,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip012",
          "item_id": "ipip012",
          "item_number": 12,
          "topic": "E3",
          "exact_item": "Take charge.",
          "key": "+",
          "domain": "E",
          "facet": "E3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Take charge."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.43,
            "probabilities": {
              "r1": 0.08,
              "r3": 0.22,
              "r4": 0.15,
              "r5": 0.01,
              "r2": 0.54
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip012",
          "item_id": "ipip012",
          "item_number": 12,
          "topic": "E3",
          "exact_item": "Take charge.",
          "key": "+",
          "domain": "E",
          "facet": "E3",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Take charge."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.39,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.21,
              "r1": 0.07,
              "r2": 0.51,
              "r3": 0.21
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip012",
          "item_id": "ipip012",
          "item_number": 12,
          "topic": "E3",
          "exact_item": "Take charge.",
          "key": "+",
          "domain": "E",
          "facet": "E3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Take charge."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.44,
            "probabilities": {
              "r1": 0.08,
              "r2": 0.56,
              "r5": 0.01,
              "r3": 0.2,
              "r4": 0.15
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip012",
          "item_id": "ipip012",
          "item_number": 12,
          "topic": "E3",
          "exact_item": "Take charge.",
          "key": "+",
          "domain": "E",
          "facet": "E3",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Take charge."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.35,
            "probabilities": {
              "r1": 0.07,
              "r3": 0.26,
              "r4": 0.19,
              "r2": 0.48,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip012",
          "item_id": "ipip012",
          "item_number": 12,
          "topic": "E3",
          "exact_item": "Take charge.",
          "key": "+",
          "domain": "E",
          "facet": "E3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Take charge."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.4,
            "probabilities": {
              "r1": 0.08,
              "r4": 0.17,
              "r5": 0.01,
              "r3": 0.22,
              "r2": 0.52
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip012",
          "item_id": "ipip012",
          "item_number": 12,
          "topic": "E3",
          "exact_item": "Take charge.",
          "key": "+",
          "domain": "E",
          "facet": "E3",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Take charge."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.39,
            "probabilities": {
              "r3": 0.21,
              "r2": 0.51,
              "r4": 0.21,
              "r5": 0.0,
              "r1": 0.07
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Experience my emotions intensely.",
      "facet": "O3",
      "domain": "O",
      "key": "+",
      "id": "ipip013",
      "number": 13,
      "original_ipip300_number": 13,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip013",
          "item_id": "ipip013",
          "item_number": 13,
          "topic": "O3",
          "exact_item": "Experience my emotions intensely.",
          "key": "+",
          "domain": "O",
          "facet": "O3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Experience my emotions intensely."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.51,
            "probabilities": {
              "r1": 0.08,
              "r3": 0.19,
              "r4": 0.12,
              "r5": 0.01,
              "r2": 0.6
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip013",
          "item_id": "ipip013",
          "item_number": 13,
          "topic": "O3",
          "exact_item": "Experience my emotions intensely.",
          "key": "+",
          "domain": "O",
          "facet": "O3",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Experience my emotions intensely."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.25,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.23,
              "r2": 0.41000000000000003,
              "r1": 0.05,
              "r3": 0.31
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip013",
          "item_id": "ipip013",
          "item_number": 13,
          "topic": "O3",
          "exact_item": "Experience my emotions intensely.",
          "key": "+",
          "domain": "O",
          "facet": "O3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Experience my emotions intensely."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.41,
            "probabilities": {
              "r1": 0.07,
              "r2": 0.52,
              "r5": 0.01,
              "r3": 0.24,
              "r4": 0.16
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip013",
          "item_id": "ipip013",
          "item_number": 13,
          "topic": "O3",
          "exact_item": "Experience my emotions intensely.",
          "key": "+",
          "domain": "O",
          "facet": "O3",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Experience my emotions intensely."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.3,
            "probabilities": {
              "r4": 0.22,
              "r3": 0.29,
              "r1": 0.05,
              "r2": 0.44,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip013",
          "item_id": "ipip013",
          "item_number": 13,
          "topic": "O3",
          "exact_item": "Experience my emotions intensely.",
          "key": "+",
          "domain": "O",
          "facet": "O3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Experience my emotions intensely."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.43,
            "probabilities": {
              "r1": 0.07,
              "r4": 0.15,
              "r5": 0.01,
              "r3": 0.22999999999999998,
              "r2": 0.54
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip013",
          "item_id": "ipip013",
          "item_number": 13,
          "topic": "O3",
          "exact_item": "Experience my emotions intensely.",
          "key": "+",
          "domain": "O",
          "facet": "O3",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Experience my emotions intensely."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.3,
            "probabilities": {
              "r3": 0.29,
              "r2": 0.44,
              "r4": 0.22,
              "r5": 0.0,
              "r1": 0.05
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Love to help others.",
      "facet": "A3",
      "domain": "A",
      "key": "+",
      "id": "ipip014",
      "number": 14,
      "original_ipip300_number": 74,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip014",
          "item_id": "ipip014",
          "item_number": 14,
          "topic": "A3",
          "exact_item": "Love to help others.",
          "key": "+",
          "domain": "A",
          "facet": "A3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love to help others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.39,
            "probabilities": {
              "r1": 0.01,
              "r3": 0.37,
              "r4": 0.52,
              "r5": 0.03,
              "r2": 0.07
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip014",
          "item_id": "ipip014",
          "item_number": 14,
          "topic": "A3",
          "exact_item": "Love to help others.",
          "key": "+",
          "domain": "A",
          "facet": "A3",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love to help others."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.43,
            "probabilities": {
              "r5": 0.02,
              "r4": 0.54,
              "r2": 0.08,
              "r1": 0.01,
              "r3": 0.34
            }
          },
          "status": "valid_with_diagnostics",
          "categorical_valid": true,
          "strict_status": "invalid_probability_mass",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": null,
          "mass": 0.9900000000000001,
          "top_keys": [
            "r4"
          ],
          "warnings": [
            "nonunit_probability_mass"
          ],
          "diagnostics": {
            "nonunit_probability_mass": true,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip014",
          "item_id": "ipip014",
          "item_number": 14,
          "topic": "A3",
          "exact_item": "Love to help others.",
          "key": "+",
          "domain": "A",
          "facet": "A3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love to help others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.31,
            "probabilities": {
              "r1": 0.01,
              "r2": 0.07,
              "r5": 0.03,
              "r3": 0.44,
              "r4": 0.45
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip014",
          "item_id": "ipip014",
          "item_number": 14,
          "topic": "A3",
          "exact_item": "Love to help others.",
          "key": "+",
          "domain": "A",
          "facet": "A3",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love to help others."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.43,
            "probabilities": {
              "r1": 0.01,
              "r3": 0.35,
              "r4": 0.54,
              "r2": 0.08,
              "r5": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip014",
          "item_id": "ipip014",
          "item_number": 14,
          "topic": "A3",
          "exact_item": "Love to help others.",
          "key": "+",
          "domain": "A",
          "facet": "A3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love to help others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.34,
            "probabilities": {
              "r1": 0.01,
              "r4": 0.48,
              "r3": 0.4,
              "r5": 0.03,
              "r2": 0.08
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip014",
          "item_id": "ipip014",
          "item_number": 14,
          "topic": "A3",
          "exact_item": "Love to help others.",
          "key": "+",
          "domain": "A",
          "facet": "A3",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love to help others."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.44,
            "probabilities": {
              "r3": 0.32,
              "r2": 0.09,
              "r4": 0.56,
              "r5": 0.02,
              "r1": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Keep my promises.",
      "facet": "C3",
      "domain": "C",
      "key": "+",
      "id": "ipip015",
      "number": 15,
      "original_ipip300_number": 45,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip015",
          "item_id": "ipip015",
          "item_number": 15,
          "topic": "C3",
          "exact_item": "Keep my promises.",
          "key": "+",
          "domain": "C",
          "facet": "C3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Keep my promises."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.35,
            "probabilities": {
              "r1": 0.0,
              "r3": 0.1,
              "r4": 0.42,
              "r5": 0.48,
              "r2": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip015",
          "item_id": "ipip015",
          "item_number": 15,
          "topic": "C3",
          "exact_item": "Keep my promises.",
          "key": "+",
          "domain": "C",
          "facet": "C3",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Keep my promises."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.41,
            "probabilities": {
              "r5": 0.53,
              "r4": 0.39,
              "r2": 0.01,
              "r1": 0.0,
              "r3": 0.07
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip015",
          "item_id": "ipip015",
          "item_number": 15,
          "topic": "C3",
          "exact_item": "Keep my promises.",
          "key": "+",
          "domain": "C",
          "facet": "C3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Keep my promises."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.34,
            "probabilities": {
              "r1": 0.0,
              "r2": 0.0,
              "r5": 0.48,
              "r3": 0.1,
              "r4": 0.42
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip015",
          "item_id": "ipip015",
          "item_number": 15,
          "topic": "C3",
          "exact_item": "Keep my promises.",
          "key": "+",
          "domain": "C",
          "facet": "C3",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Keep my promises."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.45,
            "probabilities": {
              "r1": 0.01,
              "r3": 0.07,
              "r4": 0.36,
              "r2": 0.0,
              "r5": 0.56
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip015",
          "item_id": "ipip015",
          "item_number": 15,
          "topic": "C3",
          "exact_item": "Keep my promises.",
          "key": "+",
          "domain": "C",
          "facet": "C3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Keep my promises."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.38,
            "probabilities": {
              "r1": 0.0,
              "r4": 0.41,
              "r5": 0.51,
              "r3": 0.08,
              "r2": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip015",
          "item_id": "ipip015",
          "item_number": 15,
          "topic": "C3",
          "exact_item": "Keep my promises.",
          "key": "+",
          "domain": "C",
          "facet": "C3",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Keep my promises."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.48,
            "probabilities": {
              "r3": 0.06,
              "r2": 0.0,
              "r4": 0.35,
              "r5": 0.59,
              "r1": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Find it difficult to approach others.",
      "facet": "N4",
      "domain": "N",
      "key": "+",
      "id": "ipip016",
      "number": 16,
      "original_ipip300_number": 76,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip016",
          "item_id": "ipip016",
          "item_number": 16,
          "topic": "N4",
          "exact_item": "Find it difficult to approach others.",
          "key": "+",
          "domain": "N",
          "facet": "N4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Find it difficult to approach others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.39,
            "probabilities": {
              "r1": 0.51,
              "r3": 0.02,
              "r4": 0.01,
              "r5": 0.0,
              "r2": 0.46
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip016",
          "item_id": "ipip016",
          "item_number": 16,
          "topic": "N4",
          "exact_item": "Find it difficult to approach others.",
          "key": "+",
          "domain": "N",
          "facet": "N4",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Find it difficult to approach others."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.42,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.01,
              "r2": 0.53,
              "r1": 0.42,
              "r3": 0.04
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip016",
          "item_id": "ipip016",
          "item_number": 16,
          "topic": "N4",
          "exact_item": "Find it difficult to approach others.",
          "key": "+",
          "domain": "N",
          "facet": "N4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Find it difficult to approach others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.46,
            "probabilities": {
              "r1": 0.57,
              "r2": 0.4,
              "r5": 0.0,
              "r3": 0.02,
              "r4": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip016",
          "item_id": "ipip016",
          "item_number": 16,
          "topic": "N4",
          "exact_item": "Find it difficult to approach others.",
          "key": "+",
          "domain": "N",
          "facet": "N4",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Find it difficult to approach others."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.38,
            "probabilities": {
              "r1": 0.43,
              "r3": 0.05,
              "r4": 0.01,
              "r2": 0.51,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip016",
          "item_id": "ipip016",
          "item_number": 16,
          "topic": "N4",
          "exact_item": "Find it difficult to approach others.",
          "key": "+",
          "domain": "N",
          "facet": "N4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Find it difficult to approach others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.44,
            "probabilities": {
              "r1": 0.55,
              "r4": 0.01,
              "r5": 0.0,
              "r3": 0.02,
              "r2": 0.42
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip016",
          "item_id": "ipip016",
          "item_number": 16,
          "topic": "N4",
          "exact_item": "Find it difficult to approach others.",
          "key": "+",
          "domain": "N",
          "facet": "N4",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Find it difficult to approach others."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.38,
            "probabilities": {
              "r3": 0.04,
              "r2": 0.5,
              "r4": 0.01,
              "r5": 0.0,
              "r1": 0.45
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        1,
        1
      ]
    },
    {
      "text": "Am always busy.",
      "facet": "E4",
      "domain": "E",
      "key": "+",
      "id": "ipip017",
      "number": 17,
      "original_ipip300_number": 17,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip017",
          "item_id": "ipip017",
          "item_number": 17,
          "topic": "E4",
          "exact_item": "Am always busy.",
          "key": "+",
          "domain": "E",
          "facet": "E4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am always busy."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.32,
            "probabilities": {
              "r1": 0.06,
              "r3": 0.03,
              "r4": 0.43,
              "r5": 0.03,
              "r2": 0.45
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip017",
          "item_id": "ipip017",
          "item_number": 17,
          "topic": "E4",
          "exact_item": "Am always busy.",
          "key": "+",
          "domain": "E",
          "facet": "E4",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am always busy."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.3,
            "probabilities": {
              "r5": 0.04,
              "r4": 0.43,
              "r2": 0.44,
              "r1": 0.05,
              "r3": 0.04
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip017",
          "item_id": "ipip017",
          "item_number": 17,
          "topic": "E4",
          "exact_item": "Am always busy.",
          "key": "+",
          "domain": "E",
          "facet": "E4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am always busy."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.34,
            "probabilities": {
              "r1": 0.05,
              "r2": 0.42,
              "r3": 0.03,
              "r5": 0.03,
              "r4": 0.47
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip017",
          "item_id": "ipip017",
          "item_number": 17,
          "topic": "E4",
          "exact_item": "Am always busy.",
          "key": "+",
          "domain": "E",
          "facet": "E4",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am always busy."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.34,
            "probabilities": {
              "r1": 0.05,
              "r3": 0.04,
              "r4": 0.4,
              "r2": 0.48,
              "r5": 0.03
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip017",
          "item_id": "ipip017",
          "item_number": 17,
          "topic": "E4",
          "exact_item": "Am always busy.",
          "key": "+",
          "domain": "E",
          "facet": "E4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am always busy."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.38,
            "probabilities": {
              "r1": 0.06,
              "r4": 0.38,
              "r5": 0.03,
              "r3": 0.03,
              "r2": 0.5
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip017",
          "item_id": "ipip017",
          "item_number": 17,
          "topic": "E4",
          "exact_item": "Am always busy.",
          "key": "+",
          "domain": "E",
          "facet": "E4",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am always busy."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.34,
            "probabilities": {
              "r3": 0.03,
              "r2": 0.48,
              "r4": 0.41,
              "r5": 0.02,
              "r1": 0.06
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        4
      ]
    },
    {
      "text": "Prefer variety to routine.",
      "facet": "O4",
      "domain": "O",
      "key": "+",
      "id": "ipip018",
      "number": 18,
      "original_ipip300_number": 18,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip018",
          "item_id": "ipip018",
          "item_number": 18,
          "topic": "O4",
          "exact_item": "Prefer variety to routine.",
          "key": "+",
          "domain": "O",
          "facet": "O4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Prefer variety to routine."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.99,
            "probabilities": {
              "r1": 0.0,
              "r3": 0.0,
              "r4": 0.01,
              "r5": 0.99,
              "r2": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip018",
          "item_id": "ipip018",
          "item_number": 18,
          "topic": "O4",
          "exact_item": "Prefer variety to routine.",
          "key": "+",
          "domain": "O",
          "facet": "O4",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Prefer variety to routine."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.99,
            "probabilities": {
              "r5": 1.0,
              "r4": 0.0,
              "r2": 0.0,
              "r1": 0.0,
              "r3": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip018",
          "item_id": "ipip018",
          "item_number": 18,
          "topic": "O4",
          "exact_item": "Prefer variety to routine.",
          "key": "+",
          "domain": "O",
          "facet": "O4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Prefer variety to routine."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.99,
            "probabilities": {
              "r1": 0.0,
              "r2": 0.0,
              "r5": 1.0,
              "r3": 0.0,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip018",
          "item_id": "ipip018",
          "item_number": 18,
          "topic": "O4",
          "exact_item": "Prefer variety to routine.",
          "key": "+",
          "domain": "O",
          "facet": "O4",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Prefer variety to routine."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.99,
            "probabilities": {
              "r4": 0.0,
              "r3": 0.0,
              "r1": 0.0,
              "r2": 0.0,
              "r5": 1.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip018",
          "item_id": "ipip018",
          "item_number": 18,
          "topic": "O4",
          "exact_item": "Prefer variety to routine.",
          "key": "+",
          "domain": "O",
          "facet": "O4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Prefer variety to routine."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.99,
            "probabilities": {
              "r1": 0.0,
              "r4": 0.01,
              "r5": 0.99,
              "r3": 0.0,
              "r2": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip018",
          "item_id": "ipip018",
          "item_number": 18,
          "topic": "O4",
          "exact_item": "Prefer variety to routine.",
          "key": "+",
          "domain": "O",
          "facet": "O4",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Prefer variety to routine."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.99,
            "probabilities": {
              "r3": 0.0,
              "r2": 0.0,
              "r4": 0.0,
              "r5": 1.0,
              "r1": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Love a good fight.",
      "facet": "A4",
      "domain": "A",
      "key": "-",
      "id": "ipip019",
      "number": 19,
      "original_ipip300_number": 169,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip019",
          "item_id": "ipip019",
          "item_number": 19,
          "topic": "A4",
          "exact_item": "Love a good fight.",
          "key": "-",
          "domain": "A",
          "facet": "A4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love a good fight."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.72,
            "probabilities": {
              "r1": 0.78,
              "r3": 0.0,
              "r4": 0.0,
              "r5": 0.01,
              "r2": 0.21
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip019",
          "item_id": "ipip019",
          "item_number": 19,
          "topic": "A4",
          "exact_item": "Love a good fight.",
          "key": "-",
          "domain": "A",
          "facet": "A4",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love a good fight."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.53,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.01,
              "r2": 0.35,
              "r1": 0.63,
              "r3": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip019",
          "item_id": "ipip019",
          "item_number": 19,
          "topic": "A4",
          "exact_item": "Love a good fight.",
          "key": "-",
          "domain": "A",
          "facet": "A4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love a good fight."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.72,
            "probabilities": {
              "r1": 0.77,
              "r2": 0.21,
              "r5": 0.0,
              "r3": 0.01,
              "r4": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip019",
          "item_id": "ipip019",
          "item_number": 19,
          "topic": "A4",
          "exact_item": "Love a good fight.",
          "key": "-",
          "domain": "A",
          "facet": "A4",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love a good fight."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.56,
            "probabilities": {
              "r1": 0.65,
              "r3": 0.01,
              "r4": 0.01,
              "r2": 0.33,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip019",
          "item_id": "ipip019",
          "item_number": 19,
          "topic": "A4",
          "exact_item": "Love a good fight.",
          "key": "-",
          "domain": "A",
          "facet": "A4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love a good fight."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.69,
            "probabilities": {
              "r1": 0.74,
              "r4": 0.01,
              "r5": 0.01,
              "r3": 0.01,
              "r2": 0.23
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip019",
          "item_id": "ipip019",
          "item_number": 19,
          "topic": "A4",
          "exact_item": "Love a good fight.",
          "key": "-",
          "domain": "A",
          "facet": "A4",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love a good fight."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.62,
            "probabilities": {
              "r3": 0.01,
              "r2": 0.28,
              "r4": 0.01,
              "r5": 0.0,
              "r1": 0.7
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Work hard.",
      "facet": "C4",
      "domain": "C",
      "key": "+",
      "id": "ipip020",
      "number": 20,
      "original_ipip300_number": 50,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip020",
          "item_id": "ipip020",
          "item_number": 20,
          "topic": "C4",
          "exact_item": "Work hard.",
          "key": "+",
          "domain": "C",
          "facet": "C4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Work hard."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.38,
            "probabilities": {
              "r1": 0.01,
              "r3": 0.36,
              "r4": 0.5,
              "r5": 0.11,
              "r2": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip020",
          "item_id": "ipip020",
          "item_number": 20,
          "topic": "C4",
          "exact_item": "Work hard.",
          "key": "+",
          "domain": "C",
          "facet": "C4",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Work hard."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.38,
            "probabilities": {
              "r5": 0.16,
              "r4": 0.51,
              "r1": 0.01,
              "r2": 0.02,
              "r3": 0.3
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip020",
          "item_id": "ipip020",
          "item_number": 20,
          "topic": "C4",
          "exact_item": "Work hard.",
          "key": "+",
          "domain": "C",
          "facet": "C4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Work hard."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.34,
            "probabilities": {
              "r1": 0.01,
              "r2": 0.03,
              "r5": 0.11,
              "r3": 0.38,
              "r4": 0.47
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip020",
          "item_id": "ipip020",
          "item_number": 20,
          "topic": "C4",
          "exact_item": "Work hard.",
          "key": "+",
          "domain": "C",
          "facet": "C4",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Work hard."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.43,
            "probabilities": {
              "r4": 0.54,
              "r3": 0.26,
              "r1": 0.01,
              "r2": 0.02,
              "r5": 0.16
            }
          },
          "status": "valid_with_diagnostics",
          "categorical_valid": true,
          "strict_status": "invalid_probability_mass",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": null,
          "mass": 0.9900000000000001,
          "top_keys": [
            "r4"
          ],
          "warnings": [
            "nonunit_probability_mass"
          ],
          "diagnostics": {
            "nonunit_probability_mass": true,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip020",
          "item_id": "ipip020",
          "item_number": 20,
          "topic": "C4",
          "exact_item": "Work hard.",
          "key": "+",
          "domain": "C",
          "facet": "C4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Work hard."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.39,
            "probabilities": {
              "r3": 0.34,
              "r4": 0.51,
              "r1": 0.01,
              "r5": 0.12,
              "r2": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip020",
          "item_id": "ipip020",
          "item_number": 20,
          "topic": "C4",
          "exact_item": "Work hard.",
          "key": "+",
          "domain": "C",
          "facet": "C4",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Work hard."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.42,
            "probabilities": {
              "r3": 0.27,
              "r2": 0.02,
              "r4": 0.53,
              "r5": 0.17,
              "r1": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Go on binges.",
      "facet": "N5",
      "domain": "N",
      "key": "+",
      "id": "ipip021",
      "number": 21,
      "original_ipip300_number": 111,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip021",
          "item_id": "ipip021",
          "item_number": 21,
          "topic": "N5",
          "exact_item": "Go on binges.",
          "key": "+",
          "domain": "N",
          "facet": "N5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Go on binges."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.47,
            "probabilities": {
              "r1": 0.59,
              "r3": 0.04,
              "r4": 0.01,
              "r5": 0.0,
              "r2": 0.36
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip021",
          "item_id": "ipip021",
          "item_number": 21,
          "topic": "N5",
          "exact_item": "Go on binges.",
          "key": "+",
          "domain": "N",
          "facet": "N5",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Go on binges."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.34,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.01,
              "r1": 0.45,
              "r2": 0.48,
              "r3": 0.06
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip021",
          "item_id": "ipip021",
          "item_number": 21,
          "topic": "N5",
          "exact_item": "Go on binges.",
          "key": "+",
          "domain": "N",
          "facet": "N5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Go on binges."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.44,
            "probabilities": {
              "r1": 0.55,
              "r2": 0.4,
              "r3": 0.03,
              "r5": 0.0,
              "r4": 0.01
            }
          },
          "status": "valid_with_diagnostics",
          "categorical_valid": true,
          "strict_status": "invalid_probability_mass",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": null,
          "mass": 0.9900000000000001,
          "top_keys": [
            "r1"
          ],
          "warnings": [
            "nonunit_probability_mass"
          ],
          "diagnostics": {
            "nonunit_probability_mass": true,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip021",
          "item_id": "ipip021",
          "item_number": 21,
          "topic": "N5",
          "exact_item": "Go on binges.",
          "key": "+",
          "domain": "N",
          "facet": "N5",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Go on binges."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.33,
            "probabilities": {
              "r4": 0.01,
              "r3": 0.06,
              "r1": 0.47,
              "r2": 0.46,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip021",
          "item_id": "ipip021",
          "item_number": 21,
          "topic": "N5",
          "exact_item": "Go on binges.",
          "key": "+",
          "domain": "N",
          "facet": "N5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Go on binges."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.48,
            "probabilities": {
              "r1": 0.59,
              "r4": 0.01,
              "r3": 0.03,
              "r5": 0.0,
              "r2": 0.37
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip021",
          "item_id": "ipip021",
          "item_number": 21,
          "topic": "N5",
          "exact_item": "Go on binges.",
          "key": "+",
          "domain": "N",
          "facet": "N5",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Go on binges."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.35,
            "probabilities": {
              "r3": 0.06,
              "r2": 0.44,
              "r4": 0.01,
              "r5": 0.0,
              "r1": 0.49
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        1,
        1
      ]
    },
    {
      "text": "Love excitement.",
      "facet": "E5",
      "domain": "E",
      "key": "+",
      "id": "ipip022",
      "number": 22,
      "original_ipip300_number": 22,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip022",
          "item_id": "ipip022",
          "item_number": 22,
          "topic": "E5",
          "exact_item": "Love excitement.",
          "key": "+",
          "domain": "E",
          "facet": "E5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love excitement."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.62,
            "probabilities": {
              "r1": 0.0,
              "r3": 0.04,
              "r4": 0.7,
              "r5": 0.22,
              "r2": 0.04
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip022",
          "item_id": "ipip022",
          "item_number": 22,
          "topic": "E5",
          "exact_item": "Love excitement.",
          "key": "+",
          "domain": "E",
          "facet": "E5",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love excitement."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.36,
            "probabilities": {
              "r5": 0.46,
              "r4": 0.49,
              "r2": 0.03,
              "r1": 0.01,
              "r3": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip022",
          "item_id": "ipip022",
          "item_number": 22,
          "topic": "E5",
          "exact_item": "Love excitement.",
          "key": "+",
          "domain": "E",
          "facet": "E5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love excitement."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.61,
            "probabilities": {
              "r1": 0.0,
              "r2": 0.04,
              "r5": 0.23,
              "r3": 0.04,
              "r4": 0.69
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip022",
          "item_id": "ipip022",
          "item_number": 22,
          "topic": "E5",
          "exact_item": "Love excitement.",
          "key": "+",
          "domain": "E",
          "facet": "E5",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love excitement."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.44,
            "probabilities": {
              "r1": 0.01,
              "r3": 0.01,
              "r4": 0.55,
              "r2": 0.03,
              "r5": 0.39
            }
          },
          "status": "valid_with_diagnostics",
          "categorical_valid": true,
          "strict_status": "invalid_probability_mass",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": null,
          "mass": 0.9900000000000001,
          "top_keys": [
            "r4"
          ],
          "warnings": [
            "nonunit_probability_mass"
          ],
          "diagnostics": {
            "nonunit_probability_mass": true,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip022",
          "item_id": "ipip022",
          "item_number": 22,
          "topic": "E5",
          "exact_item": "Love excitement.",
          "key": "+",
          "domain": "E",
          "facet": "E5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love excitement."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.63,
            "probabilities": {
              "r1": 0.0,
              "r4": 0.71,
              "r5": 0.22,
              "r3": 0.04,
              "r2": 0.03
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip022",
          "item_id": "ipip022",
          "item_number": 22,
          "topic": "E5",
          "exact_item": "Love excitement.",
          "key": "+",
          "domain": "E",
          "facet": "E5",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love excitement."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.39,
            "probabilities": {
              "r3": 0.01,
              "r2": 0.03,
              "r4": 0.51,
              "r5": 0.44,
              "r1": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Love to read challenging material.",
      "facet": "O5",
      "domain": "O",
      "key": "+",
      "id": "ipip023",
      "number": 23,
      "original_ipip300_number": 53,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip023",
          "item_id": "ipip023",
          "item_number": 23,
          "topic": "O5",
          "exact_item": "Love to read challenging material.",
          "key": "+",
          "domain": "O",
          "facet": "O5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love to read challenging material."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.43,
            "probabilities": {
              "r1": 0.01,
              "r3": 0.34,
              "r4": 0.54,
              "r5": 0.07,
              "r2": 0.04
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip023",
          "item_id": "ipip023",
          "item_number": 23,
          "topic": "O5",
          "exact_item": "Love to read challenging material.",
          "key": "+",
          "domain": "O",
          "facet": "O5",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love to read challenging material."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.46,
            "probabilities": {
              "r5": 0.09,
              "r4": 0.58,
              "r1": 0.01,
              "r2": 0.03,
              "r3": 0.29
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 0.9999999999999999,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip023",
          "item_id": "ipip023",
          "item_number": 23,
          "topic": "O5",
          "exact_item": "Love to read challenging material.",
          "key": "+",
          "domain": "O",
          "facet": "O5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love to read challenging material."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.38,
            "probabilities": {
              "r1": 0.01,
              "r2": 0.05,
              "r3": 0.39,
              "r5": 0.05,
              "r4": 0.5
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip023",
          "item_id": "ipip023",
          "item_number": 23,
          "topic": "O5",
          "exact_item": "Love to read challenging material.",
          "key": "+",
          "domain": "O",
          "facet": "O5",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love to read challenging material."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.48,
            "probabilities": {
              "r4": 0.58,
              "r3": 0.31,
              "r1": 0.01,
              "r2": 0.04,
              "r5": 0.06
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip023",
          "item_id": "ipip023",
          "item_number": 23,
          "topic": "O5",
          "exact_item": "Love to read challenging material.",
          "key": "+",
          "domain": "O",
          "facet": "O5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love to read challenging material."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.39,
            "probabilities": {
              "r5": 0.08,
              "r4": 0.51,
              "r3": 0.35,
              "r1": 0.01,
              "r2": 0.05
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip023",
          "item_id": "ipip023",
          "item_number": 23,
          "topic": "O5",
          "exact_item": "Love to read challenging material.",
          "key": "+",
          "domain": "O",
          "facet": "O5",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love to read challenging material."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.45,
            "probabilities": {
              "r3": 0.32,
              "r2": 0.04,
              "r4": 0.56,
              "r5": 0.07,
              "r1": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Believe that I am better than others.",
      "facet": "A5",
      "domain": "A",
      "key": "-",
      "id": "ipip024",
      "number": 24,
      "original_ipip300_number": 144,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip024",
          "item_id": "ipip024",
          "item_number": 24,
          "topic": "A5",
          "exact_item": "Believe that I am better than others.",
          "key": "-",
          "domain": "A",
          "facet": "A5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe that I am better than others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.91,
            "probabilities": {
              "r1": 0.9400000000000001,
              "r3": 0.0,
              "r4": 0.0,
              "r5": 0.0,
              "r2": 0.06
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip024",
          "item_id": "ipip024",
          "item_number": 24,
          "topic": "A5",
          "exact_item": "Believe that I am better than others.",
          "key": "-",
          "domain": "A",
          "facet": "A5",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe that I am better than others."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.87,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r1": 0.9,
              "r2": 0.09,
              "r3": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip024",
          "item_id": "ipip024",
          "item_number": 24,
          "topic": "A5",
          "exact_item": "Believe that I am better than others.",
          "key": "-",
          "domain": "A",
          "facet": "A5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe that I am better than others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.88,
            "probabilities": {
              "r1": 0.9,
              "r2": 0.09,
              "r5": 0.0,
              "r3": 0.01,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip024",
          "item_id": "ipip024",
          "item_number": 24,
          "topic": "A5",
          "exact_item": "Believe that I am better than others.",
          "key": "-",
          "domain": "A",
          "facet": "A5",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe that I am better than others."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.85,
            "probabilities": {
              "r1": 0.88,
              "r3": 0.01,
              "r4": 0.0,
              "r2": 0.11,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip024",
          "item_id": "ipip024",
          "item_number": 24,
          "topic": "A5",
          "exact_item": "Believe that I am better than others.",
          "key": "-",
          "domain": "A",
          "facet": "A5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe that I am better than others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.89,
            "probabilities": {
              "r1": 0.91,
              "r4": 0.0,
              "r3": 0.01,
              "r5": 0.0,
              "r2": 0.08
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip024",
          "item_id": "ipip024",
          "item_number": 24,
          "topic": "A5",
          "exact_item": "Believe that I am better than others.",
          "key": "-",
          "domain": "A",
          "facet": "A5",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe that I am better than others."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.84,
            "probabilities": {
              "r3": 0.01,
              "r2": 0.11,
              "r4": 0.0,
              "r5": 0.0,
              "r1": 0.88
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Am always prepared.",
      "facet": "C5",
      "domain": "C",
      "key": "+",
      "id": "ipip025",
      "number": 25,
      "original_ipip300_number": 55,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip025",
          "item_id": "ipip025",
          "item_number": 25,
          "topic": "C5",
          "exact_item": "Am always prepared.",
          "key": "+",
          "domain": "C",
          "facet": "C5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am always prepared."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.47,
            "probabilities": {
              "r1": 0.27,
              "r3": 0.07,
              "r4": 0.07,
              "r5": 0.01,
              "r2": 0.58
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip025",
          "item_id": "ipip025",
          "item_number": 25,
          "topic": "C5",
          "exact_item": "Am always prepared.",
          "key": "+",
          "domain": "C",
          "facet": "C5",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am always prepared."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.57,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.07,
              "r1": 0.16,
              "r2": 0.65,
              "r3": 0.12
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip025",
          "item_id": "ipip025",
          "item_number": 25,
          "topic": "C5",
          "exact_item": "Am always prepared.",
          "key": "+",
          "domain": "C",
          "facet": "C5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am always prepared."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.49,
            "probabilities": {
              "r1": 0.23,
              "r2": 0.59,
              "r5": 0.01,
              "r3": 0.09,
              "r4": 0.08
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip025",
          "item_id": "ipip025",
          "item_number": 25,
          "topic": "C5",
          "exact_item": "Am always prepared.",
          "key": "+",
          "domain": "C",
          "facet": "C5",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am always prepared."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.51,
            "probabilities": {
              "r1": 0.13,
              "r3": 0.16,
              "r4": 0.1,
              "r2": 0.61,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip025",
          "item_id": "ipip025",
          "item_number": 25,
          "topic": "C5",
          "exact_item": "Am always prepared.",
          "key": "+",
          "domain": "C",
          "facet": "C5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am always prepared."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.5,
            "probabilities": {
              "r1": 0.21,
              "r4": 0.1,
              "r5": 0.01,
              "r3": 0.08,
              "r2": 0.6
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip025",
          "item_id": "ipip025",
          "item_number": 25,
          "topic": "C5",
          "exact_item": "Am always prepared.",
          "key": "+",
          "domain": "C",
          "facet": "C5",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am always prepared."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.57,
            "probabilities": {
              "r3": 0.12,
              "r2": 0.66,
              "r4": 0.08,
              "r5": 0.0,
              "r1": 0.14
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Panic easily.",
      "facet": "N6",
      "domain": "N",
      "key": "+",
      "id": "ipip026",
      "number": 26,
      "original_ipip300_number": 26,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip026",
          "item_id": "ipip026",
          "item_number": 26,
          "topic": "N6",
          "exact_item": "Panic easily.",
          "key": "+",
          "domain": "N",
          "facet": "N6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Panic easily."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.64,
            "probabilities": {
              "r1": 0.71,
              "r3": 0.01,
              "r4": 0.0,
              "r5": 0.0,
              "r2": 0.28
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip026",
          "item_id": "ipip026",
          "item_number": 26,
          "topic": "N6",
          "exact_item": "Panic easily.",
          "key": "+",
          "domain": "N",
          "facet": "N6",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Panic easily."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.53,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r2": 0.36,
              "r1": 0.63,
              "r3": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip026",
          "item_id": "ipip026",
          "item_number": 26,
          "topic": "N6",
          "exact_item": "Panic easily.",
          "key": "+",
          "domain": "N",
          "facet": "N6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Panic easily."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.61,
            "probabilities": {
              "r1": 0.69,
              "r2": 0.3,
              "r3": 0.01,
              "r5": 0.0,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 0.9999999999999999,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip026",
          "item_id": "ipip026",
          "item_number": 26,
          "topic": "N6",
          "exact_item": "Panic easily.",
          "key": "+",
          "domain": "N",
          "facet": "N6",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Panic easily."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.5,
            "probabilities": {
              "r1": 0.6,
              "r3": 0.01,
              "r4": 0.0,
              "r2": 0.39,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip026",
          "item_id": "ipip026",
          "item_number": 26,
          "topic": "N6",
          "exact_item": "Panic easily.",
          "key": "+",
          "domain": "N",
          "facet": "N6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Panic easily."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.65,
            "probabilities": {
              "r1": 0.72,
              "r4": 0.0,
              "r5": 0.0,
              "r3": 0.01,
              "r2": 0.27
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip026",
          "item_id": "ipip026",
          "item_number": 26,
          "topic": "N6",
          "exact_item": "Panic easily.",
          "key": "+",
          "domain": "N",
          "facet": "N6",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Panic easily."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.48,
            "probabilities": {
              "r3": 0.01,
              "r2": 0.4,
              "r4": 0.0,
              "r5": 0.0,
              "r1": 0.59
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        1,
        1
      ]
    },
    {
      "text": "Radiate joy.",
      "facet": "E6",
      "domain": "E",
      "key": "+",
      "id": "ipip027",
      "number": 27,
      "original_ipip300_number": 27,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip027",
          "item_id": "ipip027",
          "item_number": 27,
          "topic": "E6",
          "exact_item": "Radiate joy.",
          "key": "+",
          "domain": "E",
          "facet": "E6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Radiate joy."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.37,
            "probabilities": {
              "r1": 0.07,
              "r3": 0.2,
              "r4": 0.21,
              "r5": 0.02,
              "r2": 0.5
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip027",
          "item_id": "ipip027",
          "item_number": 27,
          "topic": "E6",
          "exact_item": "Radiate joy.",
          "key": "+",
          "domain": "E",
          "facet": "E6",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Radiate joy."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.32,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.23,
              "r2": 0.46,
              "r1": 0.07,
              "r3": 0.24
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip027",
          "item_id": "ipip027",
          "item_number": 27,
          "topic": "E6",
          "exact_item": "Radiate joy.",
          "key": "+",
          "domain": "E",
          "facet": "E6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Radiate joy."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.37,
            "probabilities": {
              "r1": 0.07,
              "r2": 0.5,
              "r5": 0.01,
              "r3": 0.22,
              "r4": 0.2
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip027",
          "item_id": "ipip027",
          "item_number": 27,
          "topic": "E6",
          "exact_item": "Radiate joy.",
          "key": "+",
          "domain": "E",
          "facet": "E6",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Radiate joy."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.27,
            "probabilities": {
              "r1": 0.06,
              "r3": 0.27,
              "r4": 0.24,
              "r2": 0.43,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip027",
          "item_id": "ipip027",
          "item_number": 27,
          "topic": "E6",
          "exact_item": "Radiate joy.",
          "key": "+",
          "domain": "E",
          "facet": "E6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Radiate joy."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.4,
            "probabilities": {
              "r1": 0.08,
              "r4": 0.19,
              "r5": 0.02,
              "r3": 0.2,
              "r2": 0.51
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip027",
          "item_id": "ipip027",
          "item_number": 27,
          "topic": "E6",
          "exact_item": "Radiate joy.",
          "key": "+",
          "domain": "E",
          "facet": "E6",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Radiate joy."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.32,
            "probabilities": {
              "r3": 0.23,
              "r2": 0.46,
              "r4": 0.24,
              "r5": 0.0,
              "r1": 0.07
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Tend to vote for liberal political candidates.",
      "facet": "O6",
      "domain": "O",
      "key": "+",
      "id": "ipip028",
      "number": 28,
      "original_ipip300_number": 28,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip028",
          "item_id": "ipip028",
          "item_number": 28,
          "topic": "O6",
          "exact_item": "Tend to vote for liberal political candidates.",
          "key": "+",
          "domain": "O",
          "facet": "O6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Tend to vote for liberal political candidates."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.39,
            "probabilities": {
              "r1": 0.07,
              "r3": 0.31,
              "r4": 0.1,
              "r5": 0.01,
              "r2": 0.51
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip028",
          "item_id": "ipip028",
          "item_number": 28,
          "topic": "O6",
          "exact_item": "Tend to vote for liberal political candidates.",
          "key": "+",
          "domain": "O",
          "facet": "O6",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Tend to vote for liberal political candidates."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.32,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.08,
              "r2": 0.47000000000000003,
              "r1": 0.09,
              "r3": 0.36
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip028",
          "item_id": "ipip028",
          "item_number": 28,
          "topic": "O6",
          "exact_item": "Tend to vote for liberal political candidates.",
          "key": "+",
          "domain": "O",
          "facet": "O6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Tend to vote for liberal political candidates."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.37,
            "probabilities": {
              "r1": 0.09,
              "r2": 0.49,
              "r5": 0.01,
              "r3": 0.33,
              "r4": 0.08
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip028",
          "item_id": "ipip028",
          "item_number": 28,
          "topic": "O6",
          "exact_item": "Tend to vote for liberal political candidates.",
          "key": "+",
          "domain": "O",
          "facet": "O6",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Tend to vote for liberal political candidates."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.35,
            "probabilities": {
              "r1": 0.08,
              "r3": 0.37,
              "r4": 0.06,
              "r2": 0.49,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip028",
          "item_id": "ipip028",
          "item_number": 28,
          "topic": "O6",
          "exact_item": "Tend to vote for liberal political candidates.",
          "key": "+",
          "domain": "O",
          "facet": "O6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Tend to vote for liberal political candidates."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.42,
            "probabilities": {
              "r1": 0.11,
              "r4": 0.07,
              "r3": 0.28,
              "r5": 0.01,
              "r2": 0.53
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip028",
          "item_id": "ipip028",
          "item_number": 28,
          "topic": "O6",
          "exact_item": "Tend to vote for liberal political candidates.",
          "key": "+",
          "domain": "O",
          "facet": "O6",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Tend to vote for liberal political candidates."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.3,
            "probabilities": {
              "r3": 0.4,
              "r2": 0.44,
              "r4": 0.08,
              "r5": 0.0,
              "r1": 0.08
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Sympathize with the homeless.",
      "facet": "A6",
      "domain": "A",
      "key": "+",
      "id": "ipip029",
      "number": 29,
      "original_ipip300_number": 29,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip029",
          "item_id": "ipip029",
          "item_number": 29,
          "topic": "A6",
          "exact_item": "Sympathize with the homeless.",
          "key": "+",
          "domain": "A",
          "facet": "A6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Sympathize with the homeless."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.55,
            "probabilities": {
              "r1": 0.0,
              "r3": 0.2,
              "r4": 0.64,
              "r5": 0.15,
              "r2": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip029",
          "item_id": "ipip029",
          "item_number": 29,
          "topic": "A6",
          "exact_item": "Sympathize with the homeless.",
          "key": "+",
          "domain": "A",
          "facet": "A6",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Sympathize with the homeless."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.4,
            "probabilities": {
              "r5": 0.36,
              "r4": 0.51,
              "r2": 0.01,
              "r1": 0.01,
              "r3": 0.11
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip029",
          "item_id": "ipip029",
          "item_number": 29,
          "topic": "A6",
          "exact_item": "Sympathize with the homeless.",
          "key": "+",
          "domain": "A",
          "facet": "A6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Sympathize with the homeless."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.54,
            "probabilities": {
              "r1": 0.0,
              "r2": 0.01,
              "r5": 0.15,
              "r3": 0.21,
              "r4": 0.63
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip029",
          "item_id": "ipip029",
          "item_number": 29,
          "topic": "A6",
          "exact_item": "Sympathize with the homeless.",
          "key": "+",
          "domain": "A",
          "facet": "A6",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Sympathize with the homeless."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.4,
            "probabilities": {
              "r1": 0.01,
              "r3": 0.13,
              "r4": 0.52,
              "r2": 0.01,
              "r5": 0.33
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip029",
          "item_id": "ipip029",
          "item_number": 29,
          "topic": "A6",
          "exact_item": "Sympathize with the homeless.",
          "key": "+",
          "domain": "A",
          "facet": "A6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Sympathize with the homeless."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.55,
            "probabilities": {
              "r5": 0.13,
              "r4": 0.63,
              "r3": 0.23,
              "r1": 0.0,
              "r2": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip029",
          "item_id": "ipip029",
          "item_number": 29,
          "topic": "A6",
          "exact_item": "Sympathize with the homeless.",
          "key": "+",
          "domain": "A",
          "facet": "A6",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Sympathize with the homeless."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.39,
            "probabilities": {
              "r3": 0.17,
              "r2": 0.01,
              "r4": 0.52,
              "r5": 0.29,
              "r1": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Jump into things without thinking.",
      "facet": "C6",
      "domain": "C",
      "key": "-",
      "id": "ipip030",
      "number": 30,
      "original_ipip300_number": 120,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b01-ipip030",
          "item_id": "ipip030",
          "item_number": 30,
          "topic": "C6",
          "exact_item": "Jump into things without thinking.",
          "key": "-",
          "domain": "C",
          "facet": "C6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Jump into things without thinking."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
          "request_elapsed_seconds": 1.0364481760188937,
          "usage_reference": "calls/character_design_v2-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.62,
            "probabilities": {
              "r1": 0.18,
              "r3": 0.02,
              "r4": 0.1,
              "r5": 0.01,
              "r2": 0.69
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r1-b01-ipip030",
          "item_id": "ipip030",
          "item_number": 30,
          "topic": "C6",
          "exact_item": "Jump into things without thinking.",
          "key": "-",
          "domain": "C",
          "facet": "C6",
          "condition": "character_design_v2_reverse",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r1-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Jump into things without thinking."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
          "request_elapsed_seconds": 0.7945834669517353,
          "usage_reference": "calls/character_design_v2_reverse-r1-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.62,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.1,
              "r1": 0.17,
              "r2": 0.7,
              "r3": 0.03
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b01-ipip030",
          "item_id": "ipip030",
          "item_number": 30,
          "topic": "C6",
          "exact_item": "Jump into things without thinking.",
          "key": "-",
          "domain": "C",
          "facet": "C6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Jump into things without thinking."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
          "request_elapsed_seconds": 0.8247661440400407,
          "usage_reference": "calls/character_design_v2-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.55,
            "probabilities": {
              "r1": 0.24,
              "r2": 0.64,
              "r5": 0.01,
              "r3": 0.02,
              "r4": 0.09
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r2-b01-ipip030",
          "item_id": "ipip030",
          "item_number": 30,
          "topic": "C6",
          "exact_item": "Jump into things without thinking.",
          "key": "-",
          "domain": "C",
          "facet": "C6",
          "condition": "character_design_v2_reverse",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r2-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Jump into things without thinking."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
          "request_elapsed_seconds": 0.8069353039609268,
          "usage_reference": "calls/character_design_v2_reverse-r2-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.62,
            "probabilities": {
              "r1": 0.15,
              "r3": 0.03,
              "r4": 0.12,
              "r2": 0.7,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b01-ipip030",
          "item_id": "ipip030",
          "item_number": 30,
          "topic": "C6",
          "exact_item": "Jump into things without thinking.",
          "key": "-",
          "domain": "C",
          "facet": "C6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Jump into things without thinking."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "87a8f22267e0faeb98d5e9a673adc08ce3324c2fc5539204d8370ffa82ae9740",
          "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
          "request_elapsed_seconds": 0.8168961399933323,
          "usage_reference": "calls/character_design_v2-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.56,
            "probabilities": {
              "r1": 0.23,
              "r4": 0.1,
              "r5": 0.01,
              "r3": 0.02,
              "r2": 0.64
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2_reverse-r3-b01-ipip030",
          "item_id": "ipip030",
          "item_number": 30,
          "topic": "C6",
          "exact_item": "Jump into things without thinking.",
          "key": "-",
          "domain": "C",
          "facet": "C6",
          "condition": "character_design_v2_reverse",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2_reverse-r3-b01",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Jump into things without thinking."
          },
          "offered_choices": [
            {
              "key": "r5",
              "label": "Very Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r1",
              "label": "Very Inaccurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2_reverse",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "035074c4175f9afd2805230182ceeec514b6b35483e0c8283ff94ae04d6ae9a6",
          "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
          "request_elapsed_seconds": 0.8321334390202537,
          "usage_reference": "calls/character_design_v2_reverse-r3-b01/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.58,
            "probabilities": {
              "r3": 0.03,
              "r2": 0.67,
              "r4": 0.12,
              "r5": 0.0,
              "r1": 0.18
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Fear for the worst.",
      "facet": "N1",
      "domain": "N",
      "key": "+",
      "id": "ipip031",
      "number": 31,
      "original_ipip300_number": 31,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip031",
          "item_id": "ipip031",
          "item_number": 31,
          "topic": "N1",
          "exact_item": "Fear for the worst.",
          "key": "+",
          "domain": "N",
          "facet": "N1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Fear for the worst."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.47,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.01,
              "r1": 0.57,
              "r2": 0.4,
              "r3": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip031",
          "item_id": "ipip031",
          "item_number": 31,
          "topic": "N1",
          "exact_item": "Fear for the worst.",
          "key": "+",
          "domain": "N",
          "facet": "N1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Fear for the worst."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.5,
            "probabilities": {
              "r3": 0.02,
              "r4": 0.01,
              "r5": 0.0,
              "r1": 0.6,
              "r2": 0.37
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip031",
          "item_id": "ipip031",
          "item_number": 31,
          "topic": "N1",
          "exact_item": "Fear for the worst.",
          "key": "+",
          "domain": "N",
          "facet": "N1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Fear for the worst."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.46,
            "probabilities": {
              "r5": 0.0,
              "r2": 0.4,
              "r3": 0.02,
              "r1": 0.56,
              "r4": 0.01
            }
          },
          "status": "valid_with_diagnostics",
          "categorical_valid": true,
          "strict_status": "invalid_probability_mass",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": null,
          "mass": 0.9900000000000001,
          "top_keys": [
            "r1"
          ],
          "warnings": [
            "nonunit_probability_mass"
          ],
          "diagnostics": {
            "nonunit_probability_mass": true,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        1,
        1
      ]
    },
    {
      "text": "Feel comfortable around people.",
      "facet": "E1",
      "domain": "E",
      "key": "+",
      "id": "ipip032",
      "number": 32,
      "original_ipip300_number": 62,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip032",
          "item_id": "ipip032",
          "item_number": 32,
          "topic": "E1",
          "exact_item": "Feel comfortable around people.",
          "key": "+",
          "domain": "E",
          "facet": "E1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Feel comfortable around people."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.5,
            "probabilities": {
              "r5": 0.32,
              "r4": 0.6,
              "r2": 0.01,
              "r1": 0.0,
              "r3": 0.07
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip032",
          "item_id": "ipip032",
          "item_number": 32,
          "topic": "E1",
          "exact_item": "Feel comfortable around people.",
          "key": "+",
          "domain": "E",
          "facet": "E1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Feel comfortable around people."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.5,
            "probabilities": {
              "r3": 0.07,
              "r4": 0.6,
              "r5": 0.32,
              "r1": 0.0,
              "r2": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip032",
          "item_id": "ipip032",
          "item_number": 32,
          "topic": "E1",
          "exact_item": "Feel comfortable around people.",
          "key": "+",
          "domain": "E",
          "facet": "E1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Feel comfortable around people."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.49,
            "probabilities": {
              "r5": 0.31,
              "r4": 0.59,
              "r2": 0.01,
              "r1": 0.0,
              "r3": 0.09
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Enjoy wild flights of fantasy.",
      "facet": "O1",
      "domain": "O",
      "key": "+",
      "id": "ipip033",
      "number": 33,
      "original_ipip300_number": 33,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip033",
          "item_id": "ipip033",
          "item_number": 33,
          "topic": "O1",
          "exact_item": "Enjoy wild flights of fantasy.",
          "key": "+",
          "domain": "O",
          "facet": "O1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Enjoy wild flights of fantasy."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.21,
            "probabilities": {
              "r5": 0.02,
              "r4": 0.26,
              "r2": 0.37,
              "r1": 0.06,
              "r3": 0.29
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip033",
          "item_id": "ipip033",
          "item_number": 33,
          "topic": "O1",
          "exact_item": "Enjoy wild flights of fantasy.",
          "key": "+",
          "domain": "O",
          "facet": "O1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Enjoy wild flights of fantasy."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.25,
            "probabilities": {
              "r3": 0.28,
              "r4": 0.24,
              "r5": 0.02,
              "r1": 0.07,
              "r2": 0.39
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip033",
          "item_id": "ipip033",
          "item_number": 33,
          "topic": "O1",
          "exact_item": "Enjoy wild flights of fantasy.",
          "key": "+",
          "domain": "O",
          "facet": "O1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Enjoy wild flights of fantasy."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.19,
            "probabilities": {
              "r5": 0.02,
              "r4": 0.26,
              "r2": 0.35,
              "r1": 0.06,
              "r3": 0.31
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Believe that others have good intentions.",
      "facet": "A1",
      "domain": "A",
      "key": "+",
      "id": "ipip034",
      "number": 34,
      "original_ipip300_number": 34,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip034",
          "item_id": "ipip034",
          "item_number": 34,
          "topic": "A1",
          "exact_item": "Believe that others have good intentions.",
          "key": "+",
          "domain": "A",
          "facet": "A1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe that others have good intentions."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.59,
            "probabilities": {
              "r5": 0.13,
              "r4": 0.67,
              "r1": 0.0,
              "r2": 0.02,
              "r3": 0.18
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip034",
          "item_id": "ipip034",
          "item_number": 34,
          "topic": "A1",
          "exact_item": "Believe that others have good intentions.",
          "key": "+",
          "domain": "A",
          "facet": "A1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe that others have good intentions."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.54,
            "probabilities": {
              "r1": 0.0,
              "r4": 0.64,
              "r3": 0.23,
              "r5": 0.11,
              "r2": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip034",
          "item_id": "ipip034",
          "item_number": 34,
          "topic": "A1",
          "exact_item": "Believe that others have good intentions.",
          "key": "+",
          "domain": "A",
          "facet": "A1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe that others have good intentions."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.6,
            "probabilities": {
              "r5": 0.14,
              "r2": 0.02,
              "r3": 0.16,
              "r1": 0.0,
              "r4": 0.68
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Excel in what I do.",
      "facet": "C1",
      "domain": "C",
      "key": "+",
      "id": "ipip035",
      "number": 35,
      "original_ipip300_number": 35,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip035",
          "item_id": "ipip035",
          "item_number": 35,
          "topic": "C1",
          "exact_item": "Excel in what I do.",
          "key": "+",
          "domain": "C",
          "facet": "C1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Excel in what I do."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.27,
            "probabilities": {
              "r5": 0.05,
              "r4": 0.39999999999999997,
              "r1": 0.07,
              "r2": 0.27,
              "r3": 0.21
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip035",
          "item_id": "ipip035",
          "item_number": 35,
          "topic": "C1",
          "exact_item": "Excel in what I do.",
          "key": "+",
          "domain": "C",
          "facet": "C1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Excel in what I do."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.25,
            "probabilities": {
              "r1": 0.08,
              "r4": 0.4,
              "r3": 0.2,
              "r5": 0.05,
              "r2": 0.27
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip035",
          "item_id": "ipip035",
          "item_number": 35,
          "topic": "C1",
          "exact_item": "Excel in what I do.",
          "key": "+",
          "domain": "C",
          "facet": "C1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Excel in what I do."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.31,
            "probabilities": {
              "r5": 0.06,
              "r2": 0.22,
              "r3": 0.21,
              "r1": 0.07,
              "r4": 0.44
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Get irritated easily.",
      "facet": "N2",
      "domain": "N",
      "key": "+",
      "id": "ipip036",
      "number": 36,
      "original_ipip300_number": 36,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip036",
          "item_id": "ipip036",
          "item_number": 36,
          "topic": "N2",
          "exact_item": "Get irritated easily.",
          "key": "+",
          "domain": "N",
          "facet": "N2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Get irritated easily."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.47,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.01,
              "r2": 0.57,
              "r1": 0.4,
              "r3": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip036",
          "item_id": "ipip036",
          "item_number": 36,
          "topic": "N2",
          "exact_item": "Get irritated easily.",
          "key": "+",
          "domain": "N",
          "facet": "N2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Get irritated easily."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.43,
            "probabilities": {
              "r1": 0.43,
              "r4": 0.01,
              "r3": 0.02,
              "r5": 0.0,
              "r2": 0.54
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip036",
          "item_id": "ipip036",
          "item_number": 36,
          "topic": "N2",
          "exact_item": "Get irritated easily.",
          "key": "+",
          "domain": "N",
          "facet": "N2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Get irritated easily."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.49,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.01,
              "r2": 0.6,
              "r1": 0.37,
              "r3": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Talk to a lot of different people at parties.",
      "facet": "E2",
      "domain": "E",
      "key": "+",
      "id": "ipip037",
      "number": 37,
      "original_ipip300_number": 37,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip037",
          "item_id": "ipip037",
          "item_number": 37,
          "topic": "E2",
          "exact_item": "Talk to a lot of different people at parties.",
          "key": "+",
          "domain": "E",
          "facet": "E2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Talk to a lot of different people at parties."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.62,
            "probabilities": {
              "r5": 0.01,
              "r4": 0.17,
              "r1": 0.03,
              "r2": 0.7,
              "r3": 0.09
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip037",
          "item_id": "ipip037",
          "item_number": 37,
          "topic": "E2",
          "exact_item": "Talk to a lot of different people at parties.",
          "key": "+",
          "domain": "E",
          "facet": "E2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Talk to a lot of different people at parties."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.58,
            "probabilities": {
              "r3": 0.1,
              "r4": 0.19,
              "r5": 0.01,
              "r1": 0.02,
              "r2": 0.67
            }
          },
          "status": "valid_with_diagnostics",
          "categorical_valid": true,
          "strict_status": "invalid_probability_mass",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": null,
          "mass": 0.9900000000000001,
          "top_keys": [
            "r2"
          ],
          "warnings": [
            "nonunit_probability_mass"
          ],
          "diagnostics": {
            "nonunit_probability_mass": true,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip037",
          "item_id": "ipip037",
          "item_number": 37,
          "topic": "E2",
          "exact_item": "Talk to a lot of different people at parties.",
          "key": "+",
          "domain": "E",
          "facet": "E2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Talk to a lot of different people at parties."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.51,
            "probabilities": {
              "r5": 0.01,
              "r3": 0.13,
              "r4": 0.23,
              "r1": 0.02,
              "r2": 0.61
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "See beauty in things that others might not notice.",
      "facet": "O2",
      "domain": "O",
      "key": "+",
      "id": "ipip038",
      "number": 38,
      "original_ipip300_number": 68,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip038",
          "item_id": "ipip038",
          "item_number": 38,
          "topic": "O2",
          "exact_item": "See beauty in things that others might not notice.",
          "key": "+",
          "domain": "O",
          "facet": "O2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "See beauty in things that others might not notice."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.61,
            "probabilities": {
              "r5": 0.69,
              "r4": 0.29,
              "r1": 0.0,
              "r2": 0.0,
              "r3": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 0.9999999999999999,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip038",
          "item_id": "ipip038",
          "item_number": 38,
          "topic": "O2",
          "exact_item": "See beauty in things that others might not notice.",
          "key": "+",
          "domain": "O",
          "facet": "O2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "See beauty in things that others might not notice."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.66,
            "probabilities": {
              "r3": 0.02,
              "r4": 0.25,
              "r5": 0.73,
              "r1": 0.0,
              "r2": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip038",
          "item_id": "ipip038",
          "item_number": 38,
          "topic": "O2",
          "exact_item": "See beauty in things that others might not notice.",
          "key": "+",
          "domain": "O",
          "facet": "O2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "See beauty in things that others might not notice."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.66,
            "probabilities": {
              "r5": 0.74,
              "r4": 0.25,
              "r2": 0.0,
              "r1": 0.0,
              "r3": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Cheat to get ahead.",
      "facet": "A2",
      "domain": "A",
      "key": "-",
      "id": "ipip039",
      "number": 39,
      "original_ipip300_number": 159,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip039",
          "item_id": "ipip039",
          "item_number": 39,
          "topic": "A2",
          "exact_item": "Cheat to get ahead.",
          "key": "-",
          "domain": "A",
          "facet": "A2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Cheat to get ahead."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.93,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r1": 0.95,
              "r2": 0.05,
              "r3": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip039",
          "item_id": "ipip039",
          "item_number": 39,
          "topic": "A2",
          "exact_item": "Cheat to get ahead.",
          "key": "-",
          "domain": "A",
          "facet": "A2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Cheat to get ahead."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.94,
            "probabilities": {
              "r3": 0.0,
              "r4": 0.0,
              "r5": 0.0,
              "r1": 0.96,
              "r2": 0.04
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip039",
          "item_id": "ipip039",
          "item_number": 39,
          "topic": "A2",
          "exact_item": "Cheat to get ahead.",
          "key": "-",
          "domain": "A",
          "facet": "A2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Cheat to get ahead."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.93,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r2": 0.05,
              "r1": 0.95,
              "r3": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Often forget to put things back in their proper place.",
      "facet": "C2",
      "domain": "C",
      "key": "-",
      "id": "ipip040",
      "number": 40,
      "original_ipip300_number": 160,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip040",
          "item_id": "ipip040",
          "item_number": 40,
          "topic": "C2",
          "exact_item": "Often forget to put things back in their proper place.",
          "key": "-",
          "domain": "C",
          "facet": "C2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Often forget to put things back in their proper place."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.23,
            "probabilities": {
              "r5": 0.01,
              "r4": 0.07,
              "r1": 0.17,
              "r2": 0.38,
              "r3": 0.37
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip040",
          "item_id": "ipip040",
          "item_number": 40,
          "topic": "C2",
          "exact_item": "Often forget to put things back in their proper place.",
          "key": "-",
          "domain": "C",
          "facet": "C2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Often forget to put things back in their proper place."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.28,
            "probabilities": {
              "r3": 0.44,
              "r4": 0.04,
              "r5": 0.0,
              "r1": 0.16,
              "r2": 0.36
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip040",
          "item_id": "ipip040",
          "item_number": 40,
          "topic": "C2",
          "exact_item": "Often forget to put things back in their proper place.",
          "key": "-",
          "domain": "C",
          "facet": "C2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Often forget to put things back in their proper place."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.24,
            "probabilities": {
              "r5": 0.01,
              "r4": 0.06,
              "r2": 0.39,
              "r1": 0.16,
              "r3": 0.38
            }
          },
          "status": "valid_with_diagnostics",
          "categorical_valid": true,
          "strict_status": "returned_nonwinner",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": null,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [
            "returned_nonwinner"
          ],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": true
          }
        }
      ],
      "primary_valid_contribution_range": [
        3,
        4
      ]
    },
    {
      "text": "Dislike myself.",
      "facet": "N3",
      "domain": "N",
      "key": "+",
      "id": "ipip041",
      "number": 41,
      "original_ipip300_number": 41,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip041",
          "item_id": "ipip041",
          "item_number": 41,
          "topic": "N3",
          "exact_item": "Dislike myself.",
          "key": "+",
          "domain": "N",
          "facet": "N3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Dislike myself."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.77,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r1": 0.81,
              "r2": 0.16,
              "r3": 0.02
            }
          },
          "status": "valid_with_diagnostics",
          "categorical_valid": true,
          "strict_status": "invalid_probability_mass",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": null,
          "mass": 0.9900000000000001,
          "top_keys": [
            "r1"
          ],
          "warnings": [
            "nonunit_probability_mass"
          ],
          "diagnostics": {
            "nonunit_probability_mass": true,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip041",
          "item_id": "ipip041",
          "item_number": 41,
          "topic": "N3",
          "exact_item": "Dislike myself.",
          "key": "+",
          "domain": "N",
          "facet": "N3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Dislike myself."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.78,
            "probabilities": {
              "r1": 0.83,
              "r4": 0.0,
              "r3": 0.02,
              "r5": 0.0,
              "r2": 0.15
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip041",
          "item_id": "ipip041",
          "item_number": 41,
          "topic": "N3",
          "exact_item": "Dislike myself.",
          "key": "+",
          "domain": "N",
          "facet": "N3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Dislike myself."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.78,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r2": 0.15,
              "r1": 0.83,
              "r3": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        1,
        1
      ]
    },
    {
      "text": "Try to lead others.",
      "facet": "E3",
      "domain": "E",
      "key": "+",
      "id": "ipip042",
      "number": 42,
      "original_ipip300_number": 42,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip042",
          "item_id": "ipip042",
          "item_number": 42,
          "topic": "E3",
          "exact_item": "Try to lead others.",
          "key": "+",
          "domain": "E",
          "facet": "E3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Try to lead others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.47,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.06,
              "r1": 0.11,
              "r2": 0.59,
              "r3": 0.24
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip042",
          "item_id": "ipip042",
          "item_number": 42,
          "topic": "E3",
          "exact_item": "Try to lead others.",
          "key": "+",
          "domain": "E",
          "facet": "E3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Try to lead others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.43,
            "probabilities": {
              "r5": 0.01,
              "r4": 0.07,
              "r1": 0.13,
              "r3": 0.25,
              "r2": 0.54
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip042",
          "item_id": "ipip042",
          "item_number": 42,
          "topic": "E3",
          "exact_item": "Try to lead others.",
          "key": "+",
          "domain": "E",
          "facet": "E3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Try to lead others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.43,
            "probabilities": {
              "r5": 0.01,
              "r2": 0.54,
              "r3": 0.27,
              "r1": 0.11,
              "r4": 0.07
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Feel others' emotions.",
      "facet": "O3",
      "domain": "O",
      "key": "+",
      "id": "ipip043",
      "number": 43,
      "original_ipip300_number": 43,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip043",
          "item_id": "ipip043",
          "item_number": 43,
          "topic": "O3",
          "exact_item": "Feel others' emotions.",
          "key": "+",
          "domain": "O",
          "facet": "O3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Feel others' emotions."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.63,
            "probabilities": {
              "r5": 0.09,
              "r4": 0.71,
              "r1": 0.0,
              "r2": 0.02,
              "r3": 0.18
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip043",
          "item_id": "ipip043",
          "item_number": 43,
          "topic": "O3",
          "exact_item": "Feel others' emotions.",
          "key": "+",
          "domain": "O",
          "facet": "O3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Feel others' emotions."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.6,
            "probabilities": {
              "r1": 0.01,
              "r4": 0.67,
              "r5": 0.09,
              "r3": 0.21,
              "r2": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip043",
          "item_id": "ipip043",
          "item_number": 43,
          "topic": "O3",
          "exact_item": "Feel others' emotions.",
          "key": "+",
          "domain": "O",
          "facet": "O3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Feel others' emotions."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.62,
            "probabilities": {
              "r5": 0.1,
              "r2": 0.01,
              "r3": 0.19,
              "r1": 0.0,
              "r4": 0.7
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Am concerned about others.",
      "facet": "A3",
      "domain": "A",
      "key": "+",
      "id": "ipip044",
      "number": 44,
      "original_ipip300_number": 104,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip044",
          "item_id": "ipip044",
          "item_number": 44,
          "topic": "A3",
          "exact_item": "Am concerned about others.",
          "key": "+",
          "domain": "A",
          "facet": "A3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am concerned about others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.45,
            "probabilities": {
              "r5": 0.41,
              "r4": 0.56,
              "r1": 0.0,
              "r2": 0.0,
              "r3": 0.03
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip044",
          "item_id": "ipip044",
          "item_number": 44,
          "topic": "A3",
          "exact_item": "Am concerned about others.",
          "key": "+",
          "domain": "A",
          "facet": "A3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am concerned about others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.45,
            "probabilities": {
              "r3": 0.03,
              "r4": 0.56,
              "r5": 0.41,
              "r1": 0.0,
              "r2": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip044",
          "item_id": "ipip044",
          "item_number": 44,
          "topic": "A3",
          "exact_item": "Am concerned about others.",
          "key": "+",
          "domain": "A",
          "facet": "A3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am concerned about others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.5,
            "probabilities": {
              "r5": 0.37,
              "r2": 0.0,
              "r3": 0.02,
              "r1": 0.0,
              "r4": 0.61
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Tell the truth.",
      "facet": "C3",
      "domain": "C",
      "key": "+",
      "id": "ipip045",
      "number": 45,
      "original_ipip300_number": 105,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip045",
          "item_id": "ipip045",
          "item_number": 45,
          "topic": "C3",
          "exact_item": "Tell the truth.",
          "key": "+",
          "domain": "C",
          "facet": "C3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Tell the truth."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.72,
            "probabilities": {
              "r5": 0.78,
              "r4": 0.2,
              "r2": 0.01,
              "r1": 0.0,
              "r3": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip045",
          "item_id": "ipip045",
          "item_number": 45,
          "topic": "C3",
          "exact_item": "Tell the truth.",
          "key": "+",
          "domain": "C",
          "facet": "C3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Tell the truth."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.8,
            "probabilities": {
              "r5": 0.84,
              "r4": 0.15,
              "r1": 0.0,
              "r3": 0.01,
              "r2": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip045",
          "item_id": "ipip045",
          "item_number": 45,
          "topic": "C3",
          "exact_item": "Tell the truth.",
          "key": "+",
          "domain": "C",
          "facet": "C3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Tell the truth."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.8,
            "probabilities": {
              "r5": 0.85,
              "r4": 0.14,
              "r2": 0.0,
              "r1": 0.0,
              "r3": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Am afraid to draw attention to myself.",
      "facet": "N4",
      "domain": "N",
      "key": "+",
      "id": "ipip046",
      "number": 46,
      "original_ipip300_number": 106,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip046",
          "item_id": "ipip046",
          "item_number": 46,
          "topic": "N4",
          "exact_item": "Am afraid to draw attention to myself.",
          "key": "+",
          "domain": "N",
          "facet": "N4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am afraid to draw attention to myself."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.4,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.02,
              "r2": 0.42,
              "r1": 0.53,
              "r3": 0.03
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip046",
          "item_id": "ipip046",
          "item_number": 46,
          "topic": "N4",
          "exact_item": "Am afraid to draw attention to myself.",
          "key": "+",
          "domain": "N",
          "facet": "N4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am afraid to draw attention to myself."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.42,
            "probabilities": {
              "r1": 0.54,
              "r4": 0.02,
              "r5": 0.0,
              "r3": 0.02,
              "r2": 0.42
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip046",
          "item_id": "ipip046",
          "item_number": 46,
          "topic": "N4",
          "exact_item": "Am afraid to draw attention to myself.",
          "key": "+",
          "domain": "N",
          "facet": "N4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am afraid to draw attention to myself."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.45,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.02,
              "r2": 0.4,
              "r1": 0.56,
              "r3": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        1,
        1
      ]
    },
    {
      "text": "Am always on the go.",
      "facet": "E4",
      "domain": "E",
      "key": "+",
      "id": "ipip047",
      "number": 47,
      "original_ipip300_number": 47,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip047",
          "item_id": "ipip047",
          "item_number": 47,
          "topic": "E4",
          "exact_item": "Am always on the go.",
          "key": "+",
          "domain": "E",
          "facet": "E4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am always on the go."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.6,
            "probabilities": {
              "r5": 0.07,
              "r4": 0.68,
              "r2": 0.19,
              "r1": 0.01,
              "r3": 0.04
            }
          },
          "status": "valid_with_diagnostics",
          "categorical_valid": true,
          "strict_status": "invalid_probability_mass",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": null,
          "mass": 0.9900000000000001,
          "top_keys": [
            "r4"
          ],
          "warnings": [
            "nonunit_probability_mass"
          ],
          "diagnostics": {
            "nonunit_probability_mass": true,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip047",
          "item_id": "ipip047",
          "item_number": 47,
          "topic": "E4",
          "exact_item": "Am always on the go.",
          "key": "+",
          "domain": "E",
          "facet": "E4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am always on the go."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.55,
            "probabilities": {
              "r1": 0.02,
              "r4": 0.64,
              "r3": 0.04,
              "r5": 0.06,
              "r2": 0.24
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip047",
          "item_id": "ipip047",
          "item_number": 47,
          "topic": "E4",
          "exact_item": "Am always on the go.",
          "key": "+",
          "domain": "E",
          "facet": "E4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am always on the go."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.59,
            "probabilities": {
              "r5": 0.06,
              "r4": 0.68,
              "r2": 0.22,
              "r1": 0.01,
              "r3": 0.03
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Prefer to stick with things that I know.",
      "facet": "O4",
      "domain": "O",
      "key": "-",
      "id": "ipip048",
      "number": 48,
      "original_ipip300_number": 138,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip048",
          "item_id": "ipip048",
          "item_number": 48,
          "topic": "O4",
          "exact_item": "Prefer to stick with things that I know.",
          "key": "-",
          "domain": "O",
          "facet": "O4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Prefer to stick with things that I know."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.87,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r2": 0.1,
              "r1": 0.9,
              "r3": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip048",
          "item_id": "ipip048",
          "item_number": 48,
          "topic": "O4",
          "exact_item": "Prefer to stick with things that I know.",
          "key": "-",
          "domain": "O",
          "facet": "O4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Prefer to stick with things that I know."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.88,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r1": 0.91,
              "r3": 0.0,
              "r2": 0.09
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip048",
          "item_id": "ipip048",
          "item_number": 48,
          "topic": "O4",
          "exact_item": "Prefer to stick with things that I know.",
          "key": "-",
          "domain": "O",
          "facet": "O4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Prefer to stick with things that I know."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.85,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r2": 0.12,
              "r1": 0.88,
              "r3": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Yell at people.",
      "facet": "A4",
      "domain": "A",
      "key": "-",
      "id": "ipip049",
      "number": 49,
      "original_ipip300_number": 199,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip049",
          "item_id": "ipip049",
          "item_number": 49,
          "topic": "A4",
          "exact_item": "Yell at people.",
          "key": "-",
          "domain": "A",
          "facet": "A4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Yell at people."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.79,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r2": 0.16,
              "r1": 0.84,
              "r3": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip049",
          "item_id": "ipip049",
          "item_number": 49,
          "topic": "A4",
          "exact_item": "Yell at people.",
          "key": "-",
          "domain": "A",
          "facet": "A4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Yell at people."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.78,
            "probabilities": {
              "r3": 0.0,
              "r4": 0.0,
              "r5": 0.0,
              "r1": 0.83,
              "r2": 0.17
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip049",
          "item_id": "ipip049",
          "item_number": 49,
          "topic": "A4",
          "exact_item": "Yell at people.",
          "key": "-",
          "domain": "A",
          "facet": "A4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Yell at people."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.78,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r2": 0.16,
              "r1": 0.84,
              "r3": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Do more than what's expected of me.",
      "facet": "C4",
      "domain": "C",
      "key": "+",
      "id": "ipip050",
      "number": 50,
      "original_ipip300_number": 140,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip050",
          "item_id": "ipip050",
          "item_number": 50,
          "topic": "C4",
          "exact_item": "Do more than what's expected of me.",
          "key": "+",
          "domain": "C",
          "facet": "C4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Do more than what's expected of me."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.47,
            "probabilities": {
              "r5": 0.18,
              "r4": 0.57,
              "r1": 0.01,
              "r2": 0.05,
              "r3": 0.19
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip050",
          "item_id": "ipip050",
          "item_number": 50,
          "topic": "C4",
          "exact_item": "Do more than what's expected of me.",
          "key": "+",
          "domain": "C",
          "facet": "C4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Do more than what's expected of me."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.47,
            "probabilities": {
              "r3": 0.2,
              "r4": 0.58,
              "r5": 0.16,
              "r1": 0.01,
              "r2": 0.05
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip050",
          "item_id": "ipip050",
          "item_number": 50,
          "topic": "C4",
          "exact_item": "Do more than what's expected of me.",
          "key": "+",
          "domain": "C",
          "facet": "C4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Do more than what's expected of me."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.49,
            "probabilities": {
              "r5": 0.17,
              "r3": 0.18,
              "r4": 0.59,
              "r1": 0.01,
              "r2": 0.05
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Rarely overindulge.",
      "facet": "N5",
      "domain": "N",
      "key": "-",
      "id": "ipip051",
      "number": 51,
      "original_ipip300_number": 171,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip051",
          "item_id": "ipip051",
          "item_number": 51,
          "topic": "N5",
          "exact_item": "Rarely overindulge.",
          "key": "-",
          "domain": "N",
          "facet": "N5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Rarely overindulge."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.46,
            "probabilities": {
              "r5": 0.09,
              "r4": 0.57,
              "r1": 0.0,
              "r2": 0.02,
              "r3": 0.32
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip051",
          "item_id": "ipip051",
          "item_number": 51,
          "topic": "N5",
          "exact_item": "Rarely overindulge.",
          "key": "-",
          "domain": "N",
          "facet": "N5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Rarely overindulge."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.49,
            "probabilities": {
              "r3": 0.3,
              "r4": 0.59,
              "r5": 0.09,
              "r1": 0.0,
              "r2": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip051",
          "item_id": "ipip051",
          "item_number": 51,
          "topic": "N5",
          "exact_item": "Rarely overindulge.",
          "key": "-",
          "domain": "N",
          "facet": "N5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Rarely overindulge."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.47,
            "probabilities": {
              "r5": 0.1,
              "r4": 0.58,
              "r2": 0.02,
              "r1": 0.0,
              "r3": 0.3
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Seek adventure.",
      "facet": "E5",
      "domain": "E",
      "key": "+",
      "id": "ipip052",
      "number": 52,
      "original_ipip300_number": 52,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip052",
          "item_id": "ipip052",
          "item_number": 52,
          "topic": "E5",
          "exact_item": "Seek adventure.",
          "key": "+",
          "domain": "E",
          "facet": "E5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Seek adventure."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.58,
            "probabilities": {
              "r5": 0.67,
              "r4": 0.32,
              "r1": 0.0,
              "r2": 0.0,
              "r3": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip052",
          "item_id": "ipip052",
          "item_number": 52,
          "topic": "E5",
          "exact_item": "Seek adventure.",
          "key": "+",
          "domain": "E",
          "facet": "E5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Seek adventure."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.58,
            "probabilities": {
              "r1": 0.0,
              "r4": 0.33,
              "r3": 0.01,
              "r5": 0.66,
              "r2": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip052",
          "item_id": "ipip052",
          "item_number": 52,
          "topic": "E5",
          "exact_item": "Seek adventure.",
          "key": "+",
          "domain": "E",
          "facet": "E5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Seek adventure."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.62,
            "probabilities": {
              "r5": 0.7,
              "r2": 0.0,
              "r3": 0.0,
              "r1": 0.0,
              "r4": 0.3
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Avoid philosophical discussions.",
      "facet": "O5",
      "domain": "O",
      "key": "-",
      "id": "ipip053",
      "number": 53,
      "original_ipip300_number": 203,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip053",
          "item_id": "ipip053",
          "item_number": 53,
          "topic": "O5",
          "exact_item": "Avoid philosophical discussions.",
          "key": "-",
          "domain": "O",
          "facet": "O5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Avoid philosophical discussions."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.64,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.01,
              "r1": 0.71,
              "r2": 0.25,
              "r3": 0.03
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip053",
          "item_id": "ipip053",
          "item_number": 53,
          "topic": "O5",
          "exact_item": "Avoid philosophical discussions.",
          "key": "-",
          "domain": "O",
          "facet": "O5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Avoid philosophical discussions."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.6,
            "probabilities": {
              "r3": 0.03,
              "r4": 0.01,
              "r5": 0.0,
              "r1": 0.68,
              "r2": 0.28
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip053",
          "item_id": "ipip053",
          "item_number": 53,
          "topic": "O5",
          "exact_item": "Avoid philosophical discussions.",
          "key": "-",
          "domain": "O",
          "facet": "O5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Avoid philosophical discussions."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.66,
            "probabilities": {
              "r5": 0.01,
              "r4": 0.01,
              "r2": 0.23,
              "r1": 0.72,
              "r3": 0.03
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Think highly of myself.",
      "facet": "A5",
      "domain": "A",
      "key": "-",
      "id": "ipip054",
      "number": 54,
      "original_ipip300_number": 174,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip054",
          "item_id": "ipip054",
          "item_number": 54,
          "topic": "A5",
          "exact_item": "Think highly of myself.",
          "key": "-",
          "domain": "A",
          "facet": "A5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Think highly of myself."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.47,
            "probabilities": {
              "r5": 0.01,
              "r4": 0.07,
              "r2": 0.57,
              "r1": 0.21,
              "r3": 0.14
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip054",
          "item_id": "ipip054",
          "item_number": 54,
          "topic": "A5",
          "exact_item": "Think highly of myself.",
          "key": "-",
          "domain": "A",
          "facet": "A5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Think highly of myself."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.45,
            "probabilities": {
              "r5": 0.01,
              "r4": 0.06,
              "r1": 0.23,
              "r3": 0.15,
              "r2": 0.55
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip054",
          "item_id": "ipip054",
          "item_number": 54,
          "topic": "A5",
          "exact_item": "Think highly of myself.",
          "key": "-",
          "domain": "A",
          "facet": "A5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Think highly of myself."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.46,
            "probabilities": {
              "r5": 0.01,
              "r4": 0.07,
              "r2": 0.5599999999999999,
              "r1": 0.21,
              "r3": 0.15
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 0.9999999999999999,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Carry out my plans.",
      "facet": "C5",
      "domain": "C",
      "key": "+",
      "id": "ipip055",
      "number": 55,
      "original_ipip300_number": 145,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip055",
          "item_id": "ipip055",
          "item_number": 55,
          "topic": "C5",
          "exact_item": "Carry out my plans.",
          "key": "+",
          "domain": "C",
          "facet": "C5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Carry out my plans."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.42,
            "probabilities": {
              "r5": 0.34,
              "r4": 0.52,
              "r2": 0.02,
              "r1": 0.01,
              "r3": 0.11
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip055",
          "item_id": "ipip055",
          "item_number": 55,
          "topic": "C5",
          "exact_item": "Carry out my plans.",
          "key": "+",
          "domain": "C",
          "facet": "C5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Carry out my plans."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.41,
            "probabilities": {
              "r5": 0.35,
              "r4": 0.53,
              "r1": 0.0,
              "r3": 0.1,
              "r2": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip055",
          "item_id": "ipip055",
          "item_number": 55,
          "topic": "C5",
          "exact_item": "Carry out my plans.",
          "key": "+",
          "domain": "C",
          "facet": "C5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Carry out my plans."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.43,
            "probabilities": {
              "r5": 0.33,
              "r3": 0.1,
              "r4": 0.54,
              "r1": 0.01,
              "r2": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Become overwhelmed by events.",
      "facet": "N6",
      "domain": "N",
      "key": "+",
      "id": "ipip056",
      "number": 56,
      "original_ipip300_number": 56,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip056",
          "item_id": "ipip056",
          "item_number": 56,
          "topic": "N6",
          "exact_item": "Become overwhelmed by events.",
          "key": "+",
          "domain": "N",
          "facet": "N6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Become overwhelmed by events."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.43,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.04,
              "r2": 0.54,
              "r1": 0.38,
              "r3": 0.04
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip056",
          "item_id": "ipip056",
          "item_number": 56,
          "topic": "N6",
          "exact_item": "Become overwhelmed by events.",
          "key": "+",
          "domain": "N",
          "facet": "N6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Become overwhelmed by events."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.45,
            "probabilities": {
              "r3": 0.04,
              "r4": 0.04,
              "r5": 0.0,
              "r1": 0.36,
              "r2": 0.56
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip056",
          "item_id": "ipip056",
          "item_number": 56,
          "topic": "N6",
          "exact_item": "Become overwhelmed by events.",
          "key": "+",
          "domain": "N",
          "facet": "N6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Become overwhelmed by events."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.44,
            "probabilities": {
              "r5": 0.0,
              "r2": 0.56,
              "r3": 0.04,
              "r1": 0.36,
              "r4": 0.04
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Have a lot of fun.",
      "facet": "E6",
      "domain": "E",
      "key": "+",
      "id": "ipip057",
      "number": 57,
      "original_ipip300_number": 57,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip057",
          "item_id": "ipip057",
          "item_number": 57,
          "topic": "E6",
          "exact_item": "Have a lot of fun.",
          "key": "+",
          "domain": "E",
          "facet": "E6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Have a lot of fun."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.41,
            "probabilities": {
              "r5": 0.34,
              "r4": 0.52,
              "r1": 0.01,
              "r2": 0.02,
              "r3": 0.11
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip057",
          "item_id": "ipip057",
          "item_number": 57,
          "topic": "E6",
          "exact_item": "Have a lot of fun.",
          "key": "+",
          "domain": "E",
          "facet": "E6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Have a lot of fun."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.37,
            "probabilities": {
              "r3": 0.11,
              "r4": 0.49,
              "r5": 0.37,
              "r1": 0.01,
              "r2": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip057",
          "item_id": "ipip057",
          "item_number": 57,
          "topic": "E6",
          "exact_item": "Have a lot of fun.",
          "key": "+",
          "domain": "E",
          "facet": "E6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Have a lot of fun."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.38,
            "probabilities": {
              "r5": 0.37,
              "r3": 0.1,
              "r4": 0.5,
              "r1": 0.01,
              "r2": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Believe that there is no absolute right and wrong.",
      "facet": "O6",
      "domain": "O",
      "key": "+",
      "id": "ipip058",
      "number": 58,
      "original_ipip300_number": 58,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip058",
          "item_id": "ipip058",
          "item_number": 58,
          "topic": "O6",
          "exact_item": "Believe that there is no absolute right and wrong.",
          "key": "+",
          "domain": "O",
          "facet": "O6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe that there is no absolute right and wrong."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.34,
            "probabilities": {
              "r5": 0.08,
              "r4": 0.45999999999999996,
              "r1": 0.05,
              "r2": 0.3,
              "r3": 0.11
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip058",
          "item_id": "ipip058",
          "item_number": 58,
          "topic": "O6",
          "exact_item": "Believe that there is no absolute right and wrong.",
          "key": "+",
          "domain": "O",
          "facet": "O6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe that there is no absolute right and wrong."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.31,
            "probabilities": {
              "r3": 0.1,
              "r4": 0.45,
              "r5": 0.06,
              "r1": 0.06,
              "r2": 0.33
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip058",
          "item_id": "ipip058",
          "item_number": 58,
          "topic": "O6",
          "exact_item": "Believe that there is no absolute right and wrong.",
          "key": "+",
          "domain": "O",
          "facet": "O6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe that there is no absolute right and wrong."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.34,
            "probabilities": {
              "r5": 0.07,
              "r4": 0.48,
              "r2": 0.31,
              "r1": 0.05,
              "r3": 0.09
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Feel sympathy for those who are worse off than myself.",
      "facet": "A6",
      "domain": "A",
      "key": "+",
      "id": "ipip059",
      "number": 59,
      "original_ipip300_number": 59,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip059",
          "item_id": "ipip059",
          "item_number": 59,
          "topic": "A6",
          "exact_item": "Feel sympathy for those who are worse off than myself.",
          "key": "+",
          "domain": "A",
          "facet": "A6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Feel sympathy for those who are worse off than myself."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.47,
            "probabilities": {
              "r5": 0.28,
              "r4": 0.59,
              "r2": 0.0,
              "r1": 0.0,
              "r3": 0.13
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip059",
          "item_id": "ipip059",
          "item_number": 59,
          "topic": "A6",
          "exact_item": "Feel sympathy for those who are worse off than myself.",
          "key": "+",
          "domain": "A",
          "facet": "A6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Feel sympathy for those who are worse off than myself."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.49,
            "probabilities": {
              "r3": 0.14,
              "r4": 0.6,
              "r5": 0.26,
              "r1": 0.0,
              "r2": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip059",
          "item_id": "ipip059",
          "item_number": 59,
          "topic": "A6",
          "exact_item": "Feel sympathy for those who are worse off than myself.",
          "key": "+",
          "domain": "A",
          "facet": "A6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Feel sympathy for those who are worse off than myself."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.43,
            "probabilities": {
              "r5": 0.32,
              "r4": 0.53,
              "r2": 0.01,
              "r1": 0.0,
              "r3": 0.14
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Make rash decisions.",
      "facet": "C6",
      "domain": "C",
      "key": "-",
      "id": "ipip060",
      "number": 60,
      "original_ipip300_number": 150,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b02-ipip060",
          "item_id": "ipip060",
          "item_number": 60,
          "topic": "C6",
          "exact_item": "Make rash decisions.",
          "key": "-",
          "domain": "C",
          "facet": "C6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Make rash decisions."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
          "request_elapsed_seconds": 0.8119325170991942,
          "usage_reference": "calls/character_design_v2-r1-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.42,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.02,
              "r1": 0.43,
              "r2": 0.54,
              "r3": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b02-ipip060",
          "item_id": "ipip060",
          "item_number": 60,
          "topic": "C6",
          "exact_item": "Make rash decisions.",
          "key": "-",
          "domain": "C",
          "facet": "C6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Make rash decisions."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
          "request_elapsed_seconds": 0.7912443450186402,
          "usage_reference": "calls/character_design_v2-r2-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.45,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.02,
              "r1": 0.41,
              "r3": 0.01,
              "r2": 0.56
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b02-ipip060",
          "item_id": "ipip060",
          "item_number": 60,
          "topic": "C6",
          "exact_item": "Make rash decisions.",
          "key": "-",
          "domain": "C",
          "facet": "C6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b02",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Make rash decisions."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "18261ca84b94d7cae7c53e6405dd61a80549d87aaa2f98ade5975be049437902",
          "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
          "request_elapsed_seconds": 0.8432350450893864,
          "usage_reference": "calls/character_design_v2-r3-b02/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.46,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.03,
              "r2": 0.57,
              "r1": 0.39,
              "r3": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Am afraid of many things.",
      "facet": "N1",
      "domain": "N",
      "key": "+",
      "id": "ipip061",
      "number": 61,
      "original_ipip300_number": 61,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip061",
          "item_id": "ipip061",
          "item_number": 61,
          "topic": "N1",
          "exact_item": "Am afraid of many things.",
          "key": "+",
          "domain": "N",
          "facet": "N1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am afraid of many things."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.57,
            "probabilities": {
              "r4": 0.01,
              "r1": 0.66,
              "r3": 0.01,
              "r2": 0.32,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip061",
          "item_id": "ipip061",
          "item_number": 61,
          "topic": "N1",
          "exact_item": "Am afraid of many things.",
          "key": "+",
          "domain": "N",
          "facet": "N1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am afraid of many things."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.58,
            "probabilities": {
              "r4": 0.01,
              "r3": 0.01,
              "r2": 0.32,
              "r1": 0.66,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip061",
          "item_id": "ipip061",
          "item_number": 61,
          "topic": "N1",
          "exact_item": "Am afraid of many things.",
          "key": "+",
          "domain": "N",
          "facet": "N1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am afraid of many things."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.54,
            "probabilities": {
              "r2": 0.34,
              "r5": 0.0,
              "r4": 0.01,
              "r3": 0.01,
              "r1": 0.64
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        1,
        1
      ]
    },
    {
      "text": "Avoid contacts with others.",
      "facet": "E1",
      "domain": "E",
      "key": "-",
      "id": "ipip062",
      "number": 62,
      "original_ipip300_number": 212,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip062",
          "item_id": "ipip062",
          "item_number": 62,
          "topic": "E1",
          "exact_item": "Avoid contacts with others.",
          "key": "-",
          "domain": "E",
          "facet": "E1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Avoid contacts with others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.91,
            "probabilities": {
              "r2": 0.07,
              "r1": 0.93,
              "r3": 0.0,
              "r5": 0.0,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip062",
          "item_id": "ipip062",
          "item_number": 62,
          "topic": "E1",
          "exact_item": "Avoid contacts with others.",
          "key": "-",
          "domain": "E",
          "facet": "E1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Avoid contacts with others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.92,
            "probabilities": {
              "r5": 0.0,
              "r3": 0.0,
              "r2": 0.06,
              "r1": 0.94,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip062",
          "item_id": "ipip062",
          "item_number": 62,
          "topic": "E1",
          "exact_item": "Avoid contacts with others.",
          "key": "-",
          "domain": "E",
          "facet": "E1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Avoid contacts with others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.92,
            "probabilities": {
              "r2": 0.06,
              "r5": 0.0,
              "r4": 0.0,
              "r3": 0.0,
              "r1": 0.9400000000000001
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Love to daydream.",
      "facet": "O1",
      "domain": "O",
      "key": "+",
      "id": "ipip063",
      "number": 63,
      "original_ipip300_number": 63,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip063",
          "item_id": "ipip063",
          "item_number": 63,
          "topic": "O1",
          "exact_item": "Love to daydream.",
          "key": "+",
          "domain": "O",
          "facet": "O1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love to daydream."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.38,
            "probabilities": {
              "r4": 0.51,
              "r1": 0.02,
              "r3": 0.3,
              "r2": 0.11,
              "r5": 0.06
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip063",
          "item_id": "ipip063",
          "item_number": 63,
          "topic": "O1",
          "exact_item": "Love to daydream.",
          "key": "+",
          "domain": "O",
          "facet": "O1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love to daydream."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.41,
            "probabilities": {
              "r4": 0.53,
              "r3": 0.29,
              "r2": 0.1,
              "r1": 0.02,
              "r5": 0.06
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip063",
          "item_id": "ipip063",
          "item_number": 63,
          "topic": "O1",
          "exact_item": "Love to daydream.",
          "key": "+",
          "domain": "O",
          "facet": "O1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love to daydream."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.38,
            "probabilities": {
              "r2": 0.1,
              "r5": 0.05,
              "r3": 0.32,
              "r4": 0.51,
              "r1": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Trust what people say.",
      "facet": "A1",
      "domain": "A",
      "key": "+",
      "id": "ipip064",
      "number": 64,
      "original_ipip300_number": 64,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip064",
          "item_id": "ipip064",
          "item_number": 64,
          "topic": "A1",
          "exact_item": "Trust what people say.",
          "key": "+",
          "domain": "A",
          "facet": "A1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Trust what people say."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.29,
            "probabilities": {
              "r2": 0.43,
              "r1": 0.04,
              "r3": 0.16,
              "r5": 0.02,
              "r4": 0.35
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip064",
          "item_id": "ipip064",
          "item_number": 64,
          "topic": "A1",
          "exact_item": "Trust what people say.",
          "key": "+",
          "domain": "A",
          "facet": "A1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Trust what people say."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.31,
            "probabilities": {
              "r4": 0.33,
              "r3": 0.18,
              "r2": 0.44,
              "r1": 0.03,
              "r5": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip064",
          "item_id": "ipip064",
          "item_number": 64,
          "topic": "A1",
          "exact_item": "Trust what people say.",
          "key": "+",
          "domain": "A",
          "facet": "A1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Trust what people say."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.31,
            "probabilities": {
              "r2": 0.44,
              "r5": 0.02,
              "r3": 0.18,
              "r4": 0.33,
              "r1": 0.03
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Handle tasks smoothly.",
      "facet": "C1",
      "domain": "C",
      "key": "+",
      "id": "ipip065",
      "number": 65,
      "original_ipip300_number": 65,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip065",
          "item_id": "ipip065",
          "item_number": 65,
          "topic": "C1",
          "exact_item": "Handle tasks smoothly.",
          "key": "+",
          "domain": "C",
          "facet": "C1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Handle tasks smoothly."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.42,
            "probabilities": {
              "r2": 0.04,
              "r1": 0.01,
              "r3": 0.37,
              "r5": 0.05,
              "r4": 0.53
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip065",
          "item_id": "ipip065",
          "item_number": 65,
          "topic": "C1",
          "exact_item": "Handle tasks smoothly.",
          "key": "+",
          "domain": "C",
          "facet": "C1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Handle tasks smoothly."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.37,
            "probabilities": {
              "r5": 0.05,
              "r3": 0.39,
              "r2": 0.05,
              "r1": 0.01,
              "r4": 0.5
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip065",
          "item_id": "ipip065",
          "item_number": 65,
          "topic": "C1",
          "exact_item": "Handle tasks smoothly.",
          "key": "+",
          "domain": "C",
          "facet": "C1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Handle tasks smoothly."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.39,
            "probabilities": {
              "r2": 0.04,
              "r5": 0.06,
              "r4": 0.51,
              "r3": 0.38,
              "r1": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Lose my temper.",
      "facet": "N2",
      "domain": "N",
      "key": "+",
      "id": "ipip066",
      "number": 66,
      "original_ipip300_number": 126,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip066",
          "item_id": "ipip066",
          "item_number": 66,
          "topic": "N2",
          "exact_item": "Lose my temper.",
          "key": "+",
          "domain": "N",
          "facet": "N2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Lose my temper."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.48,
            "probabilities": {
              "r2": 0.59,
              "r1": 0.35,
              "r3": 0.03,
              "r5": 0.0,
              "r4": 0.03
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip066",
          "item_id": "ipip066",
          "item_number": 66,
          "topic": "N2",
          "exact_item": "Lose my temper.",
          "key": "+",
          "domain": "N",
          "facet": "N2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Lose my temper."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.5,
            "probabilities": {
              "r5": 0.0,
              "r3": 0.04,
              "r2": 0.61,
              "r1": 0.3,
              "r4": 0.05
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip066",
          "item_id": "ipip066",
          "item_number": 66,
          "topic": "N2",
          "exact_item": "Lose my temper.",
          "key": "+",
          "domain": "N",
          "facet": "N2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Lose my temper."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.54,
            "probabilities": {
              "r2": 0.64,
              "r5": 0.0,
              "r4": 0.05,
              "r3": 0.04,
              "r1": 0.27
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Prefer to be alone.",
      "facet": "E2",
      "domain": "E",
      "key": "-",
      "id": "ipip067",
      "number": 67,
      "original_ipip300_number": 157,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip067",
          "item_id": "ipip067",
          "item_number": 67,
          "topic": "E2",
          "exact_item": "Prefer to be alone.",
          "key": "-",
          "domain": "E",
          "facet": "E2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Prefer to be alone."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.58,
            "probabilities": {
              "r2": 0.32,
              "r1": 0.66,
              "r3": 0.01,
              "r5": 0.0,
              "r4": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip067",
          "item_id": "ipip067",
          "item_number": 67,
          "topic": "E2",
          "exact_item": "Prefer to be alone.",
          "key": "-",
          "domain": "E",
          "facet": "E2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Prefer to be alone."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.58,
            "probabilities": {
              "r5": 0.0,
              "r3": 0.01,
              "r2": 0.32,
              "r1": 0.66,
              "r4": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip067",
          "item_id": "ipip067",
          "item_number": 67,
          "topic": "E2",
          "exact_item": "Prefer to be alone.",
          "key": "-",
          "domain": "E",
          "facet": "E2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Prefer to be alone."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.57,
            "probabilities": {
              "r2": 0.33,
              "r5": 0.0,
              "r4": 0.01,
              "r3": 0.01,
              "r1": 0.65
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Do not like poetry.",
      "facet": "O2",
      "domain": "O",
      "key": "-",
      "id": "ipip068",
      "number": 68,
      "original_ipip300_number": 188,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip068",
          "item_id": "ipip068",
          "item_number": 68,
          "topic": "O2",
          "exact_item": "Do not like poetry.",
          "key": "-",
          "domain": "O",
          "facet": "O2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Do not like poetry."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.58,
            "probabilities": {
              "r2": 0.16,
              "r1": 0.15,
              "r3": 0.66,
              "r5": 0.01,
              "r4": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip068",
          "item_id": "ipip068",
          "item_number": 68,
          "topic": "O2",
          "exact_item": "Do not like poetry.",
          "key": "-",
          "domain": "O",
          "facet": "O2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Do not like poetry."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.58,
            "probabilities": {
              "r4": 0.02,
              "r3": 0.67,
              "r1": 0.14,
              "r2": 0.16,
              "r5": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip068",
          "item_id": "ipip068",
          "item_number": 68,
          "topic": "O2",
          "exact_item": "Do not like poetry.",
          "key": "-",
          "domain": "O",
          "facet": "O2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Do not like poetry."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.53,
            "probabilities": {
              "r2": 0.19,
              "r5": 0.01,
              "r4": 0.02,
              "r3": 0.62,
              "r1": 0.16
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        3,
        3
      ]
    },
    {
      "text": "Take advantage of others.",
      "facet": "A2",
      "domain": "A",
      "key": "-",
      "id": "ipip069",
      "number": 69,
      "original_ipip300_number": 249,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip069",
          "item_id": "ipip069",
          "item_number": 69,
          "topic": "A2",
          "exact_item": "Take advantage of others.",
          "key": "-",
          "domain": "A",
          "facet": "A2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Take advantage of others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.98,
            "probabilities": {
              "r2": 0.01,
              "r1": 0.99,
              "r3": 0.0,
              "r5": 0.0,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip069",
          "item_id": "ipip069",
          "item_number": 69,
          "topic": "A2",
          "exact_item": "Take advantage of others.",
          "key": "-",
          "domain": "A",
          "facet": "A2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Take advantage of others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.97,
            "probabilities": {
              "r4": 0.0,
              "r3": 0.0,
              "r1": 0.98,
              "r2": 0.02,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip069",
          "item_id": "ipip069",
          "item_number": 69,
          "topic": "A2",
          "exact_item": "Take advantage of others.",
          "key": "-",
          "domain": "A",
          "facet": "A2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Take advantage of others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.97,
            "probabilities": {
              "r2": 0.02,
              "r5": 0.0,
              "r4": 0.0,
              "r3": 0.0,
              "r1": 0.98
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Leave a mess in my room.",
      "facet": "C2",
      "domain": "C",
      "key": "-",
      "id": "ipip070",
      "number": 70,
      "original_ipip300_number": 190,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip070",
          "item_id": "ipip070",
          "item_number": 70,
          "topic": "C2",
          "exact_item": "Leave a mess in my room.",
          "key": "-",
          "domain": "C",
          "facet": "C2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Leave a mess in my room."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.23,
            "probabilities": {
              "r2": 0.38,
              "r1": 0.11,
              "r3": 0.35,
              "r5": 0.01,
              "r4": 0.15
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip070",
          "item_id": "ipip070",
          "item_number": 70,
          "topic": "C2",
          "exact_item": "Leave a mess in my room.",
          "key": "-",
          "domain": "C",
          "facet": "C2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Leave a mess in my room."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.26,
            "probabilities": {
              "r4": 0.13,
              "r3": 0.34,
              "r2": 0.41,
              "r1": 0.11,
              "r5": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip070",
          "item_id": "ipip070",
          "item_number": 70,
          "topic": "C2",
          "exact_item": "Leave a mess in my room.",
          "key": "-",
          "domain": "C",
          "facet": "C2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Leave a mess in my room."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.26,
            "probabilities": {
              "r2": 0.41,
              "r5": 0.01,
              "r4": 0.12,
              "r3": 0.33,
              "r1": 0.13
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Am often down in the dumps.",
      "facet": "N3",
      "domain": "N",
      "key": "+",
      "id": "ipip071",
      "number": 71,
      "original_ipip300_number": 71,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip071",
          "item_id": "ipip071",
          "item_number": 71,
          "topic": "N3",
          "exact_item": "Am often down in the dumps.",
          "key": "+",
          "domain": "N",
          "facet": "N3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am often down in the dumps."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.45,
            "probabilities": {
              "r5": 0.0,
              "r1": 0.56,
              "r3": 0.05,
              "r4": 0.01,
              "r2": 0.38
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip071",
          "item_id": "ipip071",
          "item_number": 71,
          "topic": "N3",
          "exact_item": "Am often down in the dumps.",
          "key": "+",
          "domain": "N",
          "facet": "N3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am often down in the dumps."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.46,
            "probabilities": {
              "r4": 0.0,
              "r3": 0.04,
              "r2": 0.39,
              "r1": 0.56,
              "r5": 0.0
            }
          },
          "status": "valid_with_diagnostics",
          "categorical_valid": true,
          "strict_status": "invalid_probability_mass",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": null,
          "mass": 0.9900000000000001,
          "top_keys": [
            "r1"
          ],
          "warnings": [
            "nonunit_probability_mass"
          ],
          "diagnostics": {
            "nonunit_probability_mass": true,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip071",
          "item_id": "ipip071",
          "item_number": 71,
          "topic": "N3",
          "exact_item": "Am often down in the dumps.",
          "key": "+",
          "domain": "N",
          "facet": "N3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am often down in the dumps."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.42,
            "probabilities": {
              "r2": 0.4,
              "r5": 0.0,
              "r4": 0.01,
              "r3": 0.05,
              "r1": 0.54
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        1,
        1
      ]
    },
    {
      "text": "Take control of things.",
      "facet": "E3",
      "domain": "E",
      "key": "+",
      "id": "ipip072",
      "number": 72,
      "original_ipip300_number": 132,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip072",
          "item_id": "ipip072",
          "item_number": 72,
          "topic": "E3",
          "exact_item": "Take control of things.",
          "key": "+",
          "domain": "E",
          "facet": "E3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Take control of things."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.46,
            "probabilities": {
              "r4": 0.12,
              "r1": 0.12,
              "r3": 0.19,
              "r2": 0.5599999999999999,
              "r5": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 0.9999999999999999,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip072",
          "item_id": "ipip072",
          "item_number": 72,
          "topic": "E3",
          "exact_item": "Take control of things.",
          "key": "+",
          "domain": "E",
          "facet": "E3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Take control of things."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.45,
            "probabilities": {
              "r5": 0.01,
              "r3": 0.2,
              "r2": 0.56,
              "r1": 0.1,
              "r4": 0.13
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip072",
          "item_id": "ipip072",
          "item_number": 72,
          "topic": "E3",
          "exact_item": "Take control of things.",
          "key": "+",
          "domain": "E",
          "facet": "E3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Take control of things."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.48,
            "probabilities": {
              "r2": 0.58,
              "r5": 0.01,
              "r4": 0.12,
              "r3": 0.19,
              "r1": 0.1
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Rarely notice my emotional reactions.",
      "facet": "O3",
      "domain": "O",
      "key": "-",
      "id": "ipip073",
      "number": 73,
      "original_ipip300_number": 223,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip073",
          "item_id": "ipip073",
          "item_number": 73,
          "topic": "O3",
          "exact_item": "Rarely notice my emotional reactions.",
          "key": "-",
          "domain": "O",
          "facet": "O3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Rarely notice my emotional reactions."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.66,
            "probabilities": {
              "r4": 0.01,
              "r1": 0.72,
              "r3": 0.01,
              "r2": 0.26,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip073",
          "item_id": "ipip073",
          "item_number": 73,
          "topic": "O3",
          "exact_item": "Rarely notice my emotional reactions.",
          "key": "-",
          "domain": "O",
          "facet": "O3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Rarely notice my emotional reactions."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.63,
            "probabilities": {
              "r4": 0.01,
              "r3": 0.01,
              "r2": 0.27,
              "r1": 0.71,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip073",
          "item_id": "ipip073",
          "item_number": 73,
          "topic": "O3",
          "exact_item": "Rarely notice my emotional reactions.",
          "key": "-",
          "domain": "O",
          "facet": "O3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Rarely notice my emotional reactions."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.63,
            "probabilities": {
              "r2": 0.28,
              "r5": 0.0,
              "r3": 0.01,
              "r4": 0.01,
              "r1": 0.7
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Am indifferent to the feelings of others.",
      "facet": "A3",
      "domain": "A",
      "key": "-",
      "id": "ipip074",
      "number": 74,
      "original_ipip300_number": 194,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip074",
          "item_id": "ipip074",
          "item_number": 74,
          "topic": "A3",
          "exact_item": "Am indifferent to the feelings of others.",
          "key": "-",
          "domain": "A",
          "facet": "A3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am indifferent to the feelings of others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.99,
            "probabilities": {
              "r2": 0.01,
              "r1": 0.99,
              "r3": 0.0,
              "r5": 0.0,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip074",
          "item_id": "ipip074",
          "item_number": 74,
          "topic": "A3",
          "exact_item": "Am indifferent to the feelings of others.",
          "key": "-",
          "domain": "A",
          "facet": "A3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am indifferent to the feelings of others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.99,
            "probabilities": {
              "r4": 0.0,
              "r3": 0.0,
              "r2": 0.01,
              "r1": 0.99,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip074",
          "item_id": "ipip074",
          "item_number": 74,
          "topic": "A3",
          "exact_item": "Am indifferent to the feelings of others.",
          "key": "-",
          "domain": "A",
          "facet": "A3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am indifferent to the feelings of others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.99,
            "probabilities": {
              "r2": 0.01,
              "r5": 0.0,
              "r4": 0.0,
              "r3": 0.0,
              "r1": 0.99
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Break rules.",
      "facet": "C3",
      "domain": "C",
      "key": "-",
      "id": "ipip075",
      "number": 75,
      "original_ipip300_number": 165,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip075",
          "item_id": "ipip075",
          "item_number": 75,
          "topic": "C3",
          "exact_item": "Break rules.",
          "key": "-",
          "domain": "C",
          "facet": "C3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Break rules."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.52,
            "probabilities": {
              "r2": 0.62,
              "r1": 0.17,
              "r3": 0.11,
              "r5": 0.01,
              "r4": 0.09
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip075",
          "item_id": "ipip075",
          "item_number": 75,
          "topic": "C3",
          "exact_item": "Break rules.",
          "key": "-",
          "domain": "C",
          "facet": "C3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Break rules."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.46,
            "probabilities": {
              "r5": 0.01,
              "r3": 0.13,
              "r2": 0.57,
              "r1": 0.19,
              "r4": 0.1
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip075",
          "item_id": "ipip075",
          "item_number": 75,
          "topic": "C3",
          "exact_item": "Break rules.",
          "key": "-",
          "domain": "C",
          "facet": "C3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Break rules."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.51,
            "probabilities": {
              "r2": 0.6,
              "r5": 0.01,
              "r4": 0.1,
              "r3": 0.12,
              "r1": 0.17
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Only feel comfortable with friends.",
      "facet": "N4",
      "domain": "N",
      "key": "+",
      "id": "ipip076",
      "number": 76,
      "original_ipip300_number": 136,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip076",
          "item_id": "ipip076",
          "item_number": 76,
          "topic": "N4",
          "exact_item": "Only feel comfortable with friends.",
          "key": "+",
          "domain": "N",
          "facet": "N4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Only feel comfortable with friends."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.42,
            "probabilities": {
              "r2": 0.37,
              "r1": 0.54,
              "r3": 0.02,
              "r5": 0.01,
              "r4": 0.06
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip076",
          "item_id": "ipip076",
          "item_number": 76,
          "topic": "N4",
          "exact_item": "Only feel comfortable with friends.",
          "key": "+",
          "domain": "N",
          "facet": "N4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Only feel comfortable with friends."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.44,
            "probabilities": {
              "r4": 0.05,
              "r3": 0.02,
              "r2": 0.38,
              "r1": 0.54,
              "r5": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip076",
          "item_id": "ipip076",
          "item_number": 76,
          "topic": "N4",
          "exact_item": "Only feel comfortable with friends.",
          "key": "+",
          "domain": "N",
          "facet": "N4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Only feel comfortable with friends."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.38,
            "probabilities": {
              "r2": 0.42,
              "r5": 0.01,
              "r4": 0.06,
              "r3": 0.02,
              "r1": 0.49
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        1,
        1
      ]
    },
    {
      "text": "Do a lot in my spare time.",
      "facet": "E4",
      "domain": "E",
      "key": "+",
      "id": "ipip077",
      "number": 77,
      "original_ipip300_number": 77,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip077",
          "item_id": "ipip077",
          "item_number": 77,
          "topic": "E4",
          "exact_item": "Do a lot in my spare time.",
          "key": "+",
          "domain": "E",
          "facet": "E4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Do a lot in my spare time."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.87,
            "probabilities": {
              "r2": 0.0,
              "r1": 0.0,
              "r3": 0.01,
              "r5": 0.9,
              "r4": 0.09
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip077",
          "item_id": "ipip077",
          "item_number": 77,
          "topic": "E4",
          "exact_item": "Do a lot in my spare time.",
          "key": "+",
          "domain": "E",
          "facet": "E4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Do a lot in my spare time."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.9,
            "probabilities": {
              "r4": 0.07,
              "r3": 0.0,
              "r2": 0.0,
              "r1": 0.0,
              "r5": 0.93
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip077",
          "item_id": "ipip077",
          "item_number": 77,
          "topic": "E4",
          "exact_item": "Do a lot in my spare time.",
          "key": "+",
          "domain": "E",
          "facet": "E4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Do a lot in my spare time."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r5",
            "confidence": 0.9,
            "probabilities": {
              "r2": 0.0,
              "r5": 0.92,
              "r4": 0.07,
              "r3": 0.01,
              "r1": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r5",
          "raw_value": 5,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r5"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Dislike changes.",
      "facet": "O4",
      "domain": "O",
      "key": "-",
      "id": "ipip078",
      "number": 78,
      "original_ipip300_number": 168,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip078",
          "item_id": "ipip078",
          "item_number": 78,
          "topic": "O4",
          "exact_item": "Dislike changes.",
          "key": "-",
          "domain": "O",
          "facet": "O4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Dislike changes."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.86,
            "probabilities": {
              "r2": 0.1,
              "r1": 0.9,
              "r3": 0.0,
              "r5": 0.0,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip078",
          "item_id": "ipip078",
          "item_number": 78,
          "topic": "O4",
          "exact_item": "Dislike changes.",
          "key": "-",
          "domain": "O",
          "facet": "O4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Dislike changes."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.83,
            "probabilities": {
              "r5": 0.0,
              "r3": 0.0,
              "r2": 0.13,
              "r1": 0.87,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip078",
          "item_id": "ipip078",
          "item_number": 78,
          "topic": "O4",
          "exact_item": "Dislike changes.",
          "key": "-",
          "domain": "O",
          "facet": "O4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Dislike changes."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.89,
            "probabilities": {
              "r2": 0.08,
              "r5": 0.0,
              "r4": 0.0,
              "r3": 0.0,
              "r1": 0.92
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Insult people.",
      "facet": "A4",
      "domain": "A",
      "key": "-",
      "id": "ipip079",
      "number": 79,
      "original_ipip300_number": 229,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip079",
          "item_id": "ipip079",
          "item_number": 79,
          "topic": "A4",
          "exact_item": "Insult people.",
          "key": "-",
          "domain": "A",
          "facet": "A4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Insult people."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.94,
            "probabilities": {
              "r4": 0.0,
              "r1": 0.96,
              "r3": 0.0,
              "r2": 0.04,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip079",
          "item_id": "ipip079",
          "item_number": 79,
          "topic": "A4",
          "exact_item": "Insult people.",
          "key": "-",
          "domain": "A",
          "facet": "A4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Insult people."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.93,
            "probabilities": {
              "r5": 0.0,
              "r3": 0.0,
              "r1": 0.95,
              "r2": 0.05,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip079",
          "item_id": "ipip079",
          "item_number": 79,
          "topic": "A4",
          "exact_item": "Insult people.",
          "key": "-",
          "domain": "A",
          "facet": "A4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Insult people."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.92,
            "probabilities": {
              "r2": 0.06,
              "r5": 0.0,
              "r4": 0.0,
              "r3": 0.0,
              "r1": 0.94
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Do just enough work to get by.",
      "facet": "C4",
      "domain": "C",
      "key": "-",
      "id": "ipip080",
      "number": 80,
      "original_ipip300_number": 260,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip080",
          "item_id": "ipip080",
          "item_number": 80,
          "topic": "C4",
          "exact_item": "Do just enough work to get by.",
          "key": "-",
          "domain": "C",
          "facet": "C4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Do just enough work to get by."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.82,
            "probabilities": {
              "r2": 0.13,
              "r1": 0.87,
              "r3": 0.0,
              "r5": 0.0,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip080",
          "item_id": "ipip080",
          "item_number": 80,
          "topic": "C4",
          "exact_item": "Do just enough work to get by.",
          "key": "-",
          "domain": "C",
          "facet": "C4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Do just enough work to get by."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.81,
            "probabilities": {
              "r4": 0.0,
              "r3": 0.0,
              "r2": 0.14,
              "r1": 0.86,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip080",
          "item_id": "ipip080",
          "item_number": 80,
          "topic": "C4",
          "exact_item": "Do just enough work to get by.",
          "key": "-",
          "domain": "C",
          "facet": "C4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Do just enough work to get by."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.81,
            "probabilities": {
              "r2": 0.14,
              "r5": 0.0,
              "r4": 0.0,
              "r3": 0.01,
              "r1": 0.85
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Easily resist temptations.",
      "facet": "N5",
      "domain": "N",
      "key": "-",
      "id": "ipip081",
      "number": 81,
      "original_ipip300_number": 201,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip081",
          "item_id": "ipip081",
          "item_number": 81,
          "topic": "N5",
          "exact_item": "Easily resist temptations.",
          "key": "-",
          "domain": "N",
          "facet": "N5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Easily resist temptations."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.25,
            "probabilities": {
              "r5": 0.01,
              "r1": 0.04,
              "r3": 0.4,
              "r4": 0.21,
              "r2": 0.34
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip081",
          "item_id": "ipip081",
          "item_number": 81,
          "topic": "N5",
          "exact_item": "Easily resist temptations.",
          "key": "-",
          "domain": "N",
          "facet": "N5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Easily resist temptations."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.25,
            "probabilities": {
              "r4": 0.23,
              "r3": 0.4,
              "r1": 0.04,
              "r2": 0.32,
              "r5": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip081",
          "item_id": "ipip081",
          "item_number": 81,
          "topic": "N5",
          "exact_item": "Easily resist temptations.",
          "key": "-",
          "domain": "N",
          "facet": "N5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Easily resist temptations."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.32,
            "probabilities": {
              "r2": 0.3,
              "r5": 0.01,
              "r4": 0.2,
              "r3": 0.45,
              "r1": 0.04
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        3,
        3
      ]
    },
    {
      "text": "Enjoy being reckless.",
      "facet": "E5",
      "domain": "E",
      "key": "+",
      "id": "ipip082",
      "number": 82,
      "original_ipip300_number": 142,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip082",
          "item_id": "ipip082",
          "item_number": 82,
          "topic": "E5",
          "exact_item": "Enjoy being reckless.",
          "key": "+",
          "domain": "E",
          "facet": "E5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Enjoy being reckless."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.56,
            "probabilities": {
              "r2": 0.65,
              "r1": 0.25,
              "r3": 0.02,
              "r5": 0.01,
              "r4": 0.07
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip082",
          "item_id": "ipip082",
          "item_number": 82,
          "topic": "E5",
          "exact_item": "Enjoy being reckless.",
          "key": "+",
          "domain": "E",
          "facet": "E5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Enjoy being reckless."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.55,
            "probabilities": {
              "r4": 0.08,
              "r3": 0.02,
              "r2": 0.64,
              "r1": 0.25,
              "r5": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip082",
          "item_id": "ipip082",
          "item_number": 82,
          "topic": "E5",
          "exact_item": "Enjoy being reckless.",
          "key": "+",
          "domain": "E",
          "facet": "E5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Enjoy being reckless."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.52,
            "probabilities": {
              "r2": 0.61,
              "r5": 0.01,
              "r4": 0.07,
              "r3": 0.03,
              "r1": 0.28
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Have difficulty understanding abstract ideas.",
      "facet": "O5",
      "domain": "O",
      "key": "-",
      "id": "ipip083",
      "number": 83,
      "original_ipip300_number": 233,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip083",
          "item_id": "ipip083",
          "item_number": 83,
          "topic": "O5",
          "exact_item": "Have difficulty understanding abstract ideas.",
          "key": "-",
          "domain": "O",
          "facet": "O5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Have difficulty understanding abstract ideas."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.75,
            "probabilities": {
              "r4": 0.0,
              "r1": 0.8,
              "r3": 0.01,
              "r2": 0.19,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip083",
          "item_id": "ipip083",
          "item_number": 83,
          "topic": "O5",
          "exact_item": "Have difficulty understanding abstract ideas.",
          "key": "-",
          "domain": "O",
          "facet": "O5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Have difficulty understanding abstract ideas."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.78,
            "probabilities": {
              "r4": 0.0,
              "r3": 0.01,
              "r2": 0.16,
              "r1": 0.83,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip083",
          "item_id": "ipip083",
          "item_number": 83,
          "topic": "O5",
          "exact_item": "Have difficulty understanding abstract ideas.",
          "key": "-",
          "domain": "O",
          "facet": "O5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Have difficulty understanding abstract ideas."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.76,
            "probabilities": {
              "r2": 0.18,
              "r5": 0.0,
              "r3": 0.01,
              "r4": 0.0,
              "r1": 0.81
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Have a high opinion of myself.",
      "facet": "A5",
      "domain": "A",
      "key": "-",
      "id": "ipip084",
      "number": 84,
      "original_ipip300_number": 204,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip084",
          "item_id": "ipip084",
          "item_number": 84,
          "topic": "A5",
          "exact_item": "Have a high opinion of myself.",
          "key": "-",
          "domain": "A",
          "facet": "A5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Have a high opinion of myself."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.43,
            "probabilities": {
              "r2": 0.54,
              "r1": 0.27,
              "r3": 0.13,
              "r5": 0.01,
              "r4": 0.05
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip084",
          "item_id": "ipip084",
          "item_number": 84,
          "topic": "A5",
          "exact_item": "Have a high opinion of myself.",
          "key": "-",
          "domain": "A",
          "facet": "A5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Have a high opinion of myself."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.44,
            "probabilities": {
              "r4": 0.05,
              "r3": 0.13,
              "r2": 0.55,
              "r1": 0.26,
              "r5": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip084",
          "item_id": "ipip084",
          "item_number": 84,
          "topic": "A5",
          "exact_item": "Have a high opinion of myself.",
          "key": "-",
          "domain": "A",
          "facet": "A5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Have a high opinion of myself."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.42,
            "probabilities": {
              "r2": 0.53,
              "r5": 0.01,
              "r3": 0.14,
              "r4": 0.06,
              "r1": 0.26
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Waste my time.",
      "facet": "C5",
      "domain": "C",
      "key": "-",
      "id": "ipip085",
      "number": 85,
      "original_ipip300_number": 205,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip085",
          "item_id": "ipip085",
          "item_number": 85,
          "topic": "C5",
          "exact_item": "Waste my time.",
          "key": "-",
          "domain": "C",
          "facet": "C5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Waste my time."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.45,
            "probabilities": {
              "r5": 0.01,
              "r1": 0.55,
              "r3": 0.03,
              "r4": 0.03,
              "r2": 0.38
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip085",
          "item_id": "ipip085",
          "item_number": 85,
          "topic": "C5",
          "exact_item": "Waste my time.",
          "key": "-",
          "domain": "C",
          "facet": "C5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Waste my time."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.48,
            "probabilities": {
              "r4": 0.03,
              "r3": 0.04,
              "r2": 0.34,
              "r1": 0.58,
              "r5": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip085",
          "item_id": "ipip085",
          "item_number": 85,
          "topic": "C5",
          "exact_item": "Waste my time.",
          "key": "-",
          "domain": "C",
          "facet": "C5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Waste my time."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.47,
            "probabilities": {
              "r2": 0.36,
              "r5": 0.01,
              "r4": 0.03,
              "r3": 0.03,
              "r1": 0.57
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 0.9999999999999999,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Feel that I'm unable to deal with things.",
      "facet": "N6",
      "domain": "N",
      "key": "+",
      "id": "ipip086",
      "number": 86,
      "original_ipip300_number": 86,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip086",
          "item_id": "ipip086",
          "item_number": 86,
          "topic": "N6",
          "exact_item": "Feel that I'm unable to deal with things.",
          "key": "+",
          "domain": "N",
          "facet": "N6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Feel that I'm unable to deal with things."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.81,
            "probabilities": {
              "r2": 0.14,
              "r1": 0.85,
              "r3": 0.01,
              "r5": 0.0,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip086",
          "item_id": "ipip086",
          "item_number": 86,
          "topic": "N6",
          "exact_item": "Feel that I'm unable to deal with things.",
          "key": "+",
          "domain": "N",
          "facet": "N6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Feel that I'm unable to deal with things."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.78,
            "probabilities": {
              "r5": 0.0,
              "r3": 0.01,
              "r2": 0.16,
              "r1": 0.83,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip086",
          "item_id": "ipip086",
          "item_number": 86,
          "topic": "N6",
          "exact_item": "Feel that I'm unable to deal with things.",
          "key": "+",
          "domain": "N",
          "facet": "N6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Feel that I'm unable to deal with things."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.78,
            "probabilities": {
              "r2": 0.16,
              "r5": 0.0,
              "r3": 0.01,
              "r4": 0.0,
              "r1": 0.83
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 1,
          "strict_contribution": 1,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        1,
        1
      ]
    },
    {
      "text": "Love life.",
      "facet": "E6",
      "domain": "E",
      "key": "+",
      "id": "ipip087",
      "number": 87,
      "original_ipip300_number": 147,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip087",
          "item_id": "ipip087",
          "item_number": 87,
          "topic": "E6",
          "exact_item": "Love life.",
          "key": "+",
          "domain": "E",
          "facet": "E6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love life."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.35,
            "probabilities": {
              "r2": 0.12,
              "r1": 0.14,
              "r3": 0.48,
              "r5": 0.05,
              "r4": 0.21
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip087",
          "item_id": "ipip087",
          "item_number": 87,
          "topic": "E6",
          "exact_item": "Love life.",
          "key": "+",
          "domain": "E",
          "facet": "E6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love life."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.38,
            "probabilities": {
              "r4": 0.2,
              "r3": 0.51,
              "r2": 0.11,
              "r1": 0.12,
              "r5": 0.06
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip087",
          "item_id": "ipip087",
          "item_number": 87,
          "topic": "E6",
          "exact_item": "Love life.",
          "key": "+",
          "domain": "E",
          "facet": "E6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Love life."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.38,
            "probabilities": {
              "r2": 0.12,
              "r5": 0.07,
              "r4": 0.17,
              "r3": 0.5,
              "r1": 0.14
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        3,
        3
      ]
    },
    {
      "text": "Tend to vote for conservative political candidates.",
      "facet": "O6",
      "domain": "O",
      "key": "-",
      "id": "ipip088",
      "number": 88,
      "original_ipip300_number": 148,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip088",
          "item_id": "ipip088",
          "item_number": 88,
          "topic": "O6",
          "exact_item": "Tend to vote for conservative political candidates.",
          "key": "-",
          "domain": "O",
          "facet": "O6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Tend to vote for conservative political candidates."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.41,
            "probabilities": {
              "r2": 0.44,
              "r1": 0.53,
              "r3": 0.02,
              "r5": 0.0,
              "r4": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip088",
          "item_id": "ipip088",
          "item_number": 88,
          "topic": "O6",
          "exact_item": "Tend to vote for conservative political candidates.",
          "key": "-",
          "domain": "O",
          "facet": "O6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Tend to vote for conservative political candidates."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.4,
            "probabilities": {
              "r5": 0.0,
              "r3": 0.03,
              "r1": 0.52,
              "r2": 0.45,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip088",
          "item_id": "ipip088",
          "item_number": 88,
          "topic": "O6",
          "exact_item": "Tend to vote for conservative political candidates.",
          "key": "-",
          "domain": "O",
          "facet": "O6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Tend to vote for conservative political candidates."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.35,
            "probabilities": {
              "r2": 0.48,
              "r5": 0.0,
              "r3": 0.04,
              "r4": 0.01,
              "r1": 0.47
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        5
      ]
    },
    {
      "text": "Am not interested in other people's problems.",
      "facet": "A6",
      "domain": "A",
      "key": "-",
      "id": "ipip089",
      "number": 89,
      "original_ipip300_number": 149,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip089",
          "item_id": "ipip089",
          "item_number": 89,
          "topic": "A6",
          "exact_item": "Am not interested in other people's problems.",
          "key": "-",
          "domain": "A",
          "facet": "A6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am not interested in other people's problems."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.87,
            "probabilities": {
              "r2": 0.1,
              "r1": 0.9,
              "r3": 0.0,
              "r5": 0.0,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip089",
          "item_id": "ipip089",
          "item_number": 89,
          "topic": "A6",
          "exact_item": "Am not interested in other people's problems.",
          "key": "-",
          "domain": "A",
          "facet": "A6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am not interested in other people's problems."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.88,
            "probabilities": {
              "r5": 0.0,
              "r3": 0.0,
              "r2": 0.09,
              "r1": 0.91,
              "r4": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip089",
          "item_id": "ipip089",
          "item_number": 89,
          "topic": "A6",
          "exact_item": "Am not interested in other people's problems.",
          "key": "-",
          "domain": "A",
          "facet": "A6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am not interested in other people's problems."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.87,
            "probabilities": {
              "r2": 0.1,
              "r5": 0.0,
              "r4": 0.0,
              "r3": 0.0,
              "r1": 0.9
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Rush into things.",
      "facet": "C6",
      "domain": "C",
      "key": "-",
      "id": "ipip090",
      "number": 90,
      "original_ipip300_number": 210,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b03-ipip090",
          "item_id": "ipip090",
          "item_number": 90,
          "topic": "C6",
          "exact_item": "Rush into things.",
          "key": "-",
          "domain": "C",
          "facet": "C6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Rush into things."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
          "request_elapsed_seconds": 0.773694674950093,
          "usage_reference": "calls/character_design_v2-r1-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.51,
            "probabilities": {
              "r4": 0.09,
              "r1": 0.28,
              "r3": 0.02,
              "r2": 0.6,
              "r5": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b03-ipip090",
          "item_id": "ipip090",
          "item_number": 90,
          "topic": "C6",
          "exact_item": "Rush into things.",
          "key": "-",
          "domain": "C",
          "facet": "C6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Rush into things."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
          "request_elapsed_seconds": 0.8374924709787592,
          "usage_reference": "calls/character_design_v2-r2-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.47,
            "probabilities": {
              "r4": 0.06,
              "r3": 0.02,
              "r2": 0.57,
              "r1": 0.34,
              "r5": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b03-ipip090",
          "item_id": "ipip090",
          "item_number": 90,
          "topic": "C6",
          "exact_item": "Rush into things.",
          "key": "-",
          "domain": "C",
          "facet": "C6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b03",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Rush into things."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "de5ad8fd2a2700ebeee5b45495ca0326d8233f5259a661fadfdbee548a4ab422",
          "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
          "request_elapsed_seconds": 0.855522723053582,
          "usage_reference": "calls/character_design_v2-r3-b03/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.54,
            "probabilities": {
              "r2": 0.63,
              "r5": 0.01,
              "r4": 0.08,
              "r3": 0.02,
              "r1": 0.26
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Get stressed out easily.",
      "facet": "N1",
      "domain": "N",
      "key": "+",
      "id": "ipip091",
      "number": 91,
      "original_ipip300_number": 91,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip091",
          "item_id": "ipip091",
          "item_number": 91,
          "topic": "N1",
          "exact_item": "Get stressed out easily.",
          "key": "+",
          "domain": "N",
          "facet": "N1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Get stressed out easily."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.6,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.04,
              "r1": 0.21,
              "r3": 0.07,
              "r2": 0.68
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip091",
          "item_id": "ipip091",
          "item_number": 91,
          "topic": "N1",
          "exact_item": "Get stressed out easily.",
          "key": "+",
          "domain": "N",
          "facet": "N1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Get stressed out easily."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.59,
            "probabilities": {
              "r4": 0.04,
              "r3": 0.08,
              "r2": 0.67,
              "r5": 0.0,
              "r1": 0.21
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip091",
          "item_id": "ipip091",
          "item_number": 91,
          "topic": "N1",
          "exact_item": "Get stressed out easily.",
          "key": "+",
          "domain": "N",
          "facet": "N1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Get stressed out easily."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.62,
            "probabilities": {
              "r4": 0.04,
              "r1": 0.2,
              "r5": 0.0,
              "r2": 0.7,
              "r3": 0.06
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Keep others at a distance.",
      "facet": "E1",
      "domain": "E",
      "key": "-",
      "id": "ipip092",
      "number": 92,
      "original_ipip300_number": 272,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip092",
          "item_id": "ipip092",
          "item_number": 92,
          "topic": "E1",
          "exact_item": "Keep others at a distance.",
          "key": "-",
          "domain": "E",
          "facet": "E1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Keep others at a distance."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.7,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.01,
              "r1": 0.75,
              "r3": 0.01,
              "r2": 0.23
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip092",
          "item_id": "ipip092",
          "item_number": 92,
          "topic": "E1",
          "exact_item": "Keep others at a distance.",
          "key": "-",
          "domain": "E",
          "facet": "E1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Keep others at a distance."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.62,
            "probabilities": {
              "r4": 0.01,
              "r3": 0.01,
              "r2": 0.28,
              "r5": 0.0,
              "r1": 0.7
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip092",
          "item_id": "ipip092",
          "item_number": 92,
          "topic": "E1",
          "exact_item": "Keep others at a distance.",
          "key": "-",
          "domain": "E",
          "facet": "E1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Keep others at a distance."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.68,
            "probabilities": {
              "r4": 0.01,
              "r1": 0.74,
              "r5": 0.0,
              "r2": 0.24,
              "r3": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Like to get lost in thought.",
      "facet": "O1",
      "domain": "O",
      "key": "+",
      "id": "ipip093",
      "number": 93,
      "original_ipip300_number": 93,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip093",
          "item_id": "ipip093",
          "item_number": 93,
          "topic": "O1",
          "exact_item": "Like to get lost in thought.",
          "key": "+",
          "domain": "O",
          "facet": "O1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Like to get lost in thought."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.43,
            "probabilities": {
              "r5": 0.05,
              "r4": 0.54,
              "r1": 0.03,
              "r3": 0.24,
              "r2": 0.14
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip093",
          "item_id": "ipip093",
          "item_number": 93,
          "topic": "O1",
          "exact_item": "Like to get lost in thought.",
          "key": "+",
          "domain": "O",
          "facet": "O1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Like to get lost in thought."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.41,
            "probabilities": {
              "r4": 0.52,
              "r3": 0.26,
              "r2": 0.12,
              "r5": 0.07,
              "r1": 0.03
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip093",
          "item_id": "ipip093",
          "item_number": 93,
          "topic": "O1",
          "exact_item": "Like to get lost in thought.",
          "key": "+",
          "domain": "O",
          "facet": "O1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Like to get lost in thought."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.44,
            "probabilities": {
              "r4": 0.55,
              "r1": 0.03,
              "r5": 0.07,
              "r2": 0.12,
              "r3": 0.23
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Distrust people.",
      "facet": "A1",
      "domain": "A",
      "key": "-",
      "id": "ipip094",
      "number": 94,
      "original_ipip300_number": 184,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip094",
          "item_id": "ipip094",
          "item_number": 94,
          "topic": "A1",
          "exact_item": "Distrust people.",
          "key": "-",
          "domain": "A",
          "facet": "A1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Distrust people."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.77,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r1": 0.8200000000000001,
              "r3": 0.0,
              "r2": 0.18
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip094",
          "item_id": "ipip094",
          "item_number": 94,
          "topic": "A1",
          "exact_item": "Distrust people.",
          "key": "-",
          "domain": "A",
          "facet": "A1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Distrust people."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.76,
            "probabilities": {
              "r4": 0.0,
              "r3": 0.01,
              "r2": 0.19,
              "r5": 0.0,
              "r1": 0.8
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip094",
          "item_id": "ipip094",
          "item_number": 94,
          "topic": "A1",
          "exact_item": "Distrust people.",
          "key": "-",
          "domain": "A",
          "facet": "A1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Distrust people."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.8,
            "probabilities": {
              "r4": 0.0,
              "r1": 0.84,
              "r5": 0.0,
              "r2": 0.16,
              "r3": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Know how to get things done.",
      "facet": "C1",
      "domain": "C",
      "key": "+",
      "id": "ipip095",
      "number": 95,
      "original_ipip300_number": 155,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip095",
          "item_id": "ipip095",
          "item_number": 95,
          "topic": "C1",
          "exact_item": "Know how to get things done.",
          "key": "+",
          "domain": "C",
          "facet": "C1",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Know how to get things done."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.55,
            "probabilities": {
              "r5": 0.1,
              "r4": 0.63,
              "r1": 0.01,
              "r3": 0.24,
              "r2": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip095",
          "item_id": "ipip095",
          "item_number": 95,
          "topic": "C1",
          "exact_item": "Know how to get things done.",
          "key": "+",
          "domain": "C",
          "facet": "C1",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Know how to get things done."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.56,
            "probabilities": {
              "r4": 0.66,
              "r3": 0.22,
              "r2": 0.02,
              "r5": 0.1,
              "r1": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip095",
          "item_id": "ipip095",
          "item_number": 95,
          "topic": "C1",
          "exact_item": "Know how to get things done.",
          "key": "+",
          "domain": "C",
          "facet": "C1",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Know how to get things done."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.57,
            "probabilities": {
              "r4": 0.65,
              "r1": 0.0,
              "r5": 0.11,
              "r2": 0.02,
              "r3": 0.22
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Am not easily annoyed.",
      "facet": "N2",
      "domain": "N",
      "key": "-",
      "id": "ipip096",
      "number": 96,
      "original_ipip300_number": 216,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip096",
          "item_id": "ipip096",
          "item_number": 96,
          "topic": "N2",
          "exact_item": "Am not easily annoyed.",
          "key": "-",
          "domain": "N",
          "facet": "N2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am not easily annoyed."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.63,
            "probabilities": {
              "r5": 0.18,
              "r4": 0.7,
              "r1": 0.0,
              "r3": 0.1,
              "r2": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip096",
          "item_id": "ipip096",
          "item_number": 96,
          "topic": "N2",
          "exact_item": "Am not easily annoyed.",
          "key": "-",
          "domain": "N",
          "facet": "N2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am not easily annoyed."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.62,
            "probabilities": {
              "r4": 0.69,
              "r3": 0.1,
              "r2": 0.02,
              "r5": 0.19,
              "r1": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip096",
          "item_id": "ipip096",
          "item_number": 96,
          "topic": "N2",
          "exact_item": "Am not easily annoyed.",
          "key": "-",
          "domain": "N",
          "facet": "N2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am not easily annoyed."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.65,
            "probabilities": {
              "r4": 0.72,
              "r1": 0.0,
              "r5": 0.16,
              "r2": 0.02,
              "r3": 0.1
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Avoid crowds.",
      "facet": "E2",
      "domain": "E",
      "key": "-",
      "id": "ipip097",
      "number": 97,
      "original_ipip300_number": 247,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip097",
          "item_id": "ipip097",
          "item_number": 97,
          "topic": "E2",
          "exact_item": "Avoid crowds.",
          "key": "-",
          "domain": "E",
          "facet": "E2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Avoid crowds."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.39,
            "probabilities": {
              "r5": 0.01,
              "r4": 0.15,
              "r1": 0.23,
              "r3": 0.1,
              "r2": 0.51
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip097",
          "item_id": "ipip097",
          "item_number": 97,
          "topic": "E2",
          "exact_item": "Avoid crowds.",
          "key": "-",
          "domain": "E",
          "facet": "E2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Avoid crowds."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.37,
            "probabilities": {
              "r4": 0.18,
              "r3": 0.11,
              "r2": 0.5,
              "r5": 0.02,
              "r1": 0.19
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip097",
          "item_id": "ipip097",
          "item_number": 97,
          "topic": "E2",
          "exact_item": "Avoid crowds.",
          "key": "-",
          "domain": "E",
          "facet": "E2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Avoid crowds."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.38,
            "probabilities": {
              "r4": 0.16,
              "r1": 0.24,
              "r5": 0.01,
              "r2": 0.5,
              "r3": 0.09
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Do not enjoy going to art museums.",
      "facet": "O2",
      "domain": "O",
      "key": "-",
      "id": "ipip098",
      "number": 98,
      "original_ipip300_number": 218,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip098",
          "item_id": "ipip098",
          "item_number": 98,
          "topic": "O2",
          "exact_item": "Do not enjoy going to art museums.",
          "key": "-",
          "domain": "O",
          "facet": "O2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Do not enjoy going to art museums."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.92,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r1": 0.9400000000000001,
              "r3": 0.0,
              "r2": 0.06
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip098",
          "item_id": "ipip098",
          "item_number": 98,
          "topic": "O2",
          "exact_item": "Do not enjoy going to art museums.",
          "key": "-",
          "domain": "O",
          "facet": "O2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Do not enjoy going to art museums."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.9,
            "probabilities": {
              "r4": 0.0,
              "r3": 0.0,
              "r2": 0.07,
              "r5": 0.0,
              "r1": 0.93
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip098",
          "item_id": "ipip098",
          "item_number": 98,
          "topic": "O2",
          "exact_item": "Do not enjoy going to art museums.",
          "key": "-",
          "domain": "O",
          "facet": "O2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Do not enjoy going to art museums."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.89,
            "probabilities": {
              "r4": 0.0,
              "r1": 0.92,
              "r5": 0.0,
              "r2": 0.08,
              "r3": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Obstruct others' plans.",
      "facet": "A2",
      "domain": "A",
      "key": "-",
      "id": "ipip099",
      "number": 99,
      "original_ipip300_number": 279,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip099",
          "item_id": "ipip099",
          "item_number": 99,
          "topic": "A2",
          "exact_item": "Obstruct others' plans.",
          "key": "-",
          "domain": "A",
          "facet": "A2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Obstruct others' plans."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.93,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r1": 0.95,
              "r3": 0.0,
              "r2": 0.05
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip099",
          "item_id": "ipip099",
          "item_number": 99,
          "topic": "A2",
          "exact_item": "Obstruct others' plans.",
          "key": "-",
          "domain": "A",
          "facet": "A2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Obstruct others' plans."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.94,
            "probabilities": {
              "r4": 0.0,
              "r3": 0.0,
              "r2": 0.04,
              "r5": 0.0,
              "r1": 0.96
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip099",
          "item_id": "ipip099",
          "item_number": 99,
          "topic": "A2",
          "exact_item": "Obstruct others' plans.",
          "key": "-",
          "domain": "A",
          "facet": "A2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Obstruct others' plans."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.94,
            "probabilities": {
              "r4": 0.0,
              "r1": 0.96,
              "r5": 0.0,
              "r2": 0.04,
              "r3": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Leave my belongings around.",
      "facet": "C2",
      "domain": "C",
      "key": "-",
      "id": "ipip100",
      "number": 100,
      "original_ipip300_number": 220,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip100",
          "item_id": "ipip100",
          "item_number": 100,
          "topic": "C2",
          "exact_item": "Leave my belongings around.",
          "key": "-",
          "domain": "C",
          "facet": "C2",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Leave my belongings around."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.28,
            "probabilities": {
              "r5": 0.01,
              "r4": 0.17,
              "r2": 0.33,
              "r3": 0.43,
              "r1": 0.06
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip100",
          "item_id": "ipip100",
          "item_number": 100,
          "topic": "C2",
          "exact_item": "Leave my belongings around.",
          "key": "-",
          "domain": "C",
          "facet": "C2",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Leave my belongings around."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.3,
            "probabilities": {
              "r4": 0.17,
              "r3": 0.44,
              "r2": 0.32,
              "r5": 0.01,
              "r1": 0.06
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip100",
          "item_id": "ipip100",
          "item_number": 100,
          "topic": "C2",
          "exact_item": "Leave my belongings around.",
          "key": "-",
          "domain": "C",
          "facet": "C2",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Leave my belongings around."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.28,
            "probabilities": {
              "r4": 0.15,
              "r1": 0.06,
              "r5": 0.01,
              "r2": 0.36,
              "r3": 0.42
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        3,
        3
      ]
    },
    {
      "text": "Feel comfortable with myself.",
      "facet": "N3",
      "domain": "N",
      "key": "-",
      "id": "ipip101",
      "number": 101,
      "original_ipip300_number": 251,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip101",
          "item_id": "ipip101",
          "item_number": 101,
          "topic": "N3",
          "exact_item": "Feel comfortable with myself.",
          "key": "-",
          "domain": "N",
          "facet": "N3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Feel comfortable with myself."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.42,
            "probabilities": {
              "r5": 0.27,
              "r4": 0.53,
              "r1": 0.0,
              "r3": 0.19,
              "r2": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip101",
          "item_id": "ipip101",
          "item_number": 101,
          "topic": "N3",
          "exact_item": "Feel comfortable with myself.",
          "key": "-",
          "domain": "N",
          "facet": "N3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Feel comfortable with myself."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.42,
            "probabilities": {
              "r4": 0.54,
              "r3": 0.18,
              "r2": 0.01,
              "r5": 0.27,
              "r1": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip101",
          "item_id": "ipip101",
          "item_number": 101,
          "topic": "N3",
          "exact_item": "Feel comfortable with myself.",
          "key": "-",
          "domain": "N",
          "facet": "N3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Feel comfortable with myself."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.42,
            "probabilities": {
              "r4": 0.55,
              "r1": 0.0,
              "r3": 0.19,
              "r2": 0.01,
              "r5": 0.25
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Wait for others to lead the way.",
      "facet": "E3",
      "domain": "E",
      "key": "-",
      "id": "ipip102",
      "number": 102,
      "original_ipip300_number": 162,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip102",
          "item_id": "ipip102",
          "item_number": 102,
          "topic": "E3",
          "exact_item": "Wait for others to lead the way.",
          "key": "-",
          "domain": "E",
          "facet": "E3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Wait for others to lead the way."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.43,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.01,
              "r1": 0.55,
              "r3": 0.01,
              "r2": 0.43
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip102",
          "item_id": "ipip102",
          "item_number": 102,
          "topic": "E3",
          "exact_item": "Wait for others to lead the way.",
          "key": "-",
          "domain": "E",
          "facet": "E3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Wait for others to lead the way."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.49,
            "probabilities": {
              "r4": 0.01,
              "r3": 0.01,
              "r2": 0.39,
              "r5": 0.0,
              "r1": 0.59
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip102",
          "item_id": "ipip102",
          "item_number": 102,
          "topic": "E3",
          "exact_item": "Wait for others to lead the way.",
          "key": "-",
          "domain": "E",
          "facet": "E3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Wait for others to lead the way."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.43,
            "probabilities": {
              "r4": 0.01,
              "r1": 0.55,
              "r5": 0.0,
              "r2": 0.43,
              "r3": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Don't understand people who get emotional.",
      "facet": "O3",
      "domain": "O",
      "key": "-",
      "id": "ipip103",
      "number": 103,
      "original_ipip300_number": 283,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip103",
          "item_id": "ipip103",
          "item_number": 103,
          "topic": "O3",
          "exact_item": "Don't understand people who get emotional.",
          "key": "-",
          "domain": "O",
          "facet": "O3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Don't understand people who get emotional."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.91,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r1": 0.9400000000000001,
              "r3": 0.0,
              "r2": 0.06
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip103",
          "item_id": "ipip103",
          "item_number": 103,
          "topic": "O3",
          "exact_item": "Don't understand people who get emotional.",
          "key": "-",
          "domain": "O",
          "facet": "O3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Don't understand people who get emotional."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.9,
            "probabilities": {
              "r4": 0.0,
              "r3": 0.0,
              "r2": 0.07,
              "r5": 0.0,
              "r1": 0.93
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip103",
          "item_id": "ipip103",
          "item_number": 103,
          "topic": "O3",
          "exact_item": "Don't understand people who get emotional.",
          "key": "-",
          "domain": "O",
          "facet": "O3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Don't understand people who get emotional."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.9,
            "probabilities": {
              "r4": 0.0,
              "r1": 0.92,
              "r3": 0.0,
              "r2": 0.08,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Take no time for others.",
      "facet": "A3",
      "domain": "A",
      "key": "-",
      "id": "ipip104",
      "number": 104,
      "original_ipip300_number": 284,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip104",
          "item_id": "ipip104",
          "item_number": 104,
          "topic": "A3",
          "exact_item": "Take no time for others.",
          "key": "-",
          "domain": "A",
          "facet": "A3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Take no time for others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.94,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r1": 0.95,
              "r3": 0.0,
              "r2": 0.05
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip104",
          "item_id": "ipip104",
          "item_number": 104,
          "topic": "A3",
          "exact_item": "Take no time for others.",
          "key": "-",
          "domain": "A",
          "facet": "A3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Take no time for others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.93,
            "probabilities": {
              "r4": 0.0,
              "r3": 0.0,
              "r2": 0.05,
              "r5": 0.0,
              "r1": 0.95
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip104",
          "item_id": "ipip104",
          "item_number": 104,
          "topic": "A3",
          "exact_item": "Take no time for others.",
          "key": "-",
          "domain": "A",
          "facet": "A3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Take no time for others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.94,
            "probabilities": {
              "r4": 0.0,
              "r1": 0.96,
              "r5": 0.0,
              "r2": 0.04,
              "r3": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Break my promises.",
      "facet": "C3",
      "domain": "C",
      "key": "-",
      "id": "ipip105",
      "number": 105,
      "original_ipip300_number": 195,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip105",
          "item_id": "ipip105",
          "item_number": 105,
          "topic": "C3",
          "exact_item": "Break my promises.",
          "key": "-",
          "domain": "C",
          "facet": "C3",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Break my promises."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.86,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r1": 0.89,
              "r3": 0.0,
              "r2": 0.11
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip105",
          "item_id": "ipip105",
          "item_number": 105,
          "topic": "C3",
          "exact_item": "Break my promises.",
          "key": "-",
          "domain": "C",
          "facet": "C3",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Break my promises."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.81,
            "probabilities": {
              "r4": 0.0,
              "r3": 0.01,
              "r2": 0.14,
              "r5": 0.0,
              "r1": 0.85
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip105",
          "item_id": "ipip105",
          "item_number": 105,
          "topic": "C3",
          "exact_item": "Break my promises.",
          "key": "-",
          "domain": "C",
          "facet": "C3",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Break my promises."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.84,
            "probabilities": {
              "r4": 0.0,
              "r1": 0.88,
              "r5": 0.0,
              "r2": 0.12,
              "r3": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Am not bothered by difficult social situations.",
      "facet": "N4",
      "domain": "N",
      "key": "-",
      "id": "ipip106",
      "number": 106,
      "original_ipip300_number": 256,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip106",
          "item_id": "ipip106",
          "item_number": 106,
          "topic": "N4",
          "exact_item": "Am not bothered by difficult social situations.",
          "key": "-",
          "domain": "N",
          "facet": "N4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am not bothered by difficult social situations."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.36,
            "probabilities": {
              "r5": 0.1,
              "r4": 0.5,
              "r1": 0.03,
              "r3": 0.15,
              "r2": 0.22
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip106",
          "item_id": "ipip106",
          "item_number": 106,
          "topic": "N4",
          "exact_item": "Am not bothered by difficult social situations.",
          "key": "-",
          "domain": "N",
          "facet": "N4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am not bothered by difficult social situations."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.4,
            "probabilities": {
              "r4": 0.52,
              "r3": 0.15,
              "r2": 0.21,
              "r5": 0.1,
              "r1": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip106",
          "item_id": "ipip106",
          "item_number": 106,
          "topic": "N4",
          "exact_item": "Am not bothered by difficult social situations.",
          "key": "-",
          "domain": "N",
          "facet": "N4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am not bothered by difficult social situations."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.37,
            "probabilities": {
              "r4": 0.5,
              "r1": 0.03,
              "r3": 0.14,
              "r2": 0.24,
              "r5": 0.09
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Like to take it easy.",
      "facet": "E4",
      "domain": "E",
      "key": "-",
      "id": "ipip107",
      "number": 107,
      "original_ipip300_number": 167,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip107",
          "item_id": "ipip107",
          "item_number": 107,
          "topic": "E4",
          "exact_item": "Like to take it easy.",
          "key": "-",
          "domain": "E",
          "facet": "E4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Like to take it easy."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.53,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.01,
              "r1": 0.63,
              "r3": 0.0,
              "r2": 0.36
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip107",
          "item_id": "ipip107",
          "item_number": 107,
          "topic": "E4",
          "exact_item": "Like to take it easy.",
          "key": "-",
          "domain": "E",
          "facet": "E4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Like to take it easy."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.55,
            "probabilities": {
              "r4": 0.01,
              "r3": 0.0,
              "r2": 0.34,
              "r5": 0.0,
              "r1": 0.65
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip107",
          "item_id": "ipip107",
          "item_number": 107,
          "topic": "E4",
          "exact_item": "Like to take it easy.",
          "key": "-",
          "domain": "E",
          "facet": "E4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Like to take it easy."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.55,
            "probabilities": {
              "r4": 0.01,
              "r1": 0.65,
              "r5": 0.0,
              "r2": 0.34,
              "r3": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Am attached to conventional ways.",
      "facet": "O4",
      "domain": "O",
      "key": "-",
      "id": "ipip108",
      "number": 108,
      "original_ipip300_number": 288,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip108",
          "item_id": "ipip108",
          "item_number": 108,
          "topic": "O4",
          "exact_item": "Am attached to conventional ways.",
          "key": "-",
          "domain": "O",
          "facet": "O4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am attached to conventional ways."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.38,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.01,
              "r1": 0.5,
              "r3": 0.01,
              "r2": 0.48
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip108",
          "item_id": "ipip108",
          "item_number": 108,
          "topic": "O4",
          "exact_item": "Am attached to conventional ways.",
          "key": "-",
          "domain": "O",
          "facet": "O4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am attached to conventional ways."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.36,
            "probabilities": {
              "r4": 0.01,
              "r3": 0.01,
              "r2": 0.49,
              "r5": 0.0,
              "r1": 0.49
            }
          },
          "status": "valid_with_diagnostics",
          "categorical_valid": true,
          "strict_status": "tied_top",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": null,
          "mass": 1.0,
          "top_keys": [
            "r2",
            "r1"
          ],
          "warnings": [
            "tied_top"
          ],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": true,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip108",
          "item_id": "ipip108",
          "item_number": 108,
          "topic": "O4",
          "exact_item": "Am attached to conventional ways.",
          "key": "-",
          "domain": "O",
          "facet": "O4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am attached to conventional ways."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.37,
            "probabilities": {
              "r4": 0.01,
              "r1": 0.48,
              "r3": 0.01,
              "r2": 0.5,
              "r5": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        5
      ]
    },
    {
      "text": "Get back at others.",
      "facet": "A4",
      "domain": "A",
      "key": "-",
      "id": "ipip109",
      "number": 109,
      "original_ipip300_number": 259,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip109",
          "item_id": "ipip109",
          "item_number": 109,
          "topic": "A4",
          "exact_item": "Get back at others.",
          "key": "-",
          "domain": "A",
          "facet": "A4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Get back at others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.88,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r1": 0.91,
              "r3": 0.0,
              "r2": 0.09
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip109",
          "item_id": "ipip109",
          "item_number": 109,
          "topic": "A4",
          "exact_item": "Get back at others.",
          "key": "-",
          "domain": "A",
          "facet": "A4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Get back at others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.88,
            "probabilities": {
              "r4": 0.0,
              "r3": 0.0,
              "r2": 0.09,
              "r5": 0.0,
              "r1": 0.91
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip109",
          "item_id": "ipip109",
          "item_number": 109,
          "topic": "A4",
          "exact_item": "Get back at others.",
          "key": "-",
          "domain": "A",
          "facet": "A4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Get back at others."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.88,
            "probabilities": {
              "r4": 0.0,
              "r1": 0.91,
              "r5": 0.0,
              "r2": 0.09,
              "r3": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Put little time and effort into my work.",
      "facet": "C4",
      "domain": "C",
      "key": "-",
      "id": "ipip110",
      "number": 110,
      "original_ipip300_number": 290,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip110",
          "item_id": "ipip110",
          "item_number": 110,
          "topic": "C4",
          "exact_item": "Put little time and effort into my work.",
          "key": "-",
          "domain": "C",
          "facet": "C4",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Put little time and effort into my work."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.94,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r2": 0.04,
              "r3": 0.0,
              "r1": 0.96
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip110",
          "item_id": "ipip110",
          "item_number": 110,
          "topic": "C4",
          "exact_item": "Put little time and effort into my work.",
          "key": "-",
          "domain": "C",
          "facet": "C4",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Put little time and effort into my work."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.95,
            "probabilities": {
              "r4": 0.0,
              "r3": 0.0,
              "r2": 0.04,
              "r5": 0.0,
              "r1": 0.96
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip110",
          "item_id": "ipip110",
          "item_number": 110,
          "topic": "C4",
          "exact_item": "Put little time and effort into my work.",
          "key": "-",
          "domain": "C",
          "facet": "C4",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Put little time and effort into my work."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.95,
            "probabilities": {
              "r4": 0.0,
              "r1": 0.97,
              "r5": 0.0,
              "r2": 0.03,
              "r3": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Am able to control my cravings.",
      "facet": "N5",
      "domain": "N",
      "key": "-",
      "id": "ipip111",
      "number": 111,
      "original_ipip300_number": 231,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip111",
          "item_id": "ipip111",
          "item_number": 111,
          "topic": "N5",
          "exact_item": "Am able to control my cravings.",
          "key": "-",
          "domain": "N",
          "facet": "N5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am able to control my cravings."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.54,
            "probabilities": {
              "r5": 0.02,
              "r4": 0.31,
              "r1": 0.01,
              "r3": 0.64,
              "r2": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip111",
          "item_id": "ipip111",
          "item_number": 111,
          "topic": "N5",
          "exact_item": "Am able to control my cravings.",
          "key": "-",
          "domain": "N",
          "facet": "N5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am able to control my cravings."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.61,
            "probabilities": {
              "r4": 0.27,
              "r3": 0.69,
              "r2": 0.01,
              "r5": 0.02,
              "r1": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip111",
          "item_id": "ipip111",
          "item_number": 111,
          "topic": "N5",
          "exact_item": "Am able to control my cravings.",
          "key": "-",
          "domain": "N",
          "facet": "N5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am able to control my cravings."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Neither Inaccurate nor Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r3",
            "confidence": 0.56,
            "probabilities": {
              "r4": 0.31,
              "r1": 0.01,
              "r5": 0.02,
              "r2": 0.02,
              "r3": 0.64
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r3",
          "raw_value": 3,
          "contribution": 3,
          "strict_contribution": 3,
          "mass": 1.0,
          "top_keys": [
            "r3"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        3,
        3
      ]
    },
    {
      "text": "Act wild and crazy.",
      "facet": "E5",
      "domain": "E",
      "key": "+",
      "id": "ipip112",
      "number": 112,
      "original_ipip300_number": 172,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip112",
          "item_id": "ipip112",
          "item_number": 112,
          "topic": "E5",
          "exact_item": "Act wild and crazy.",
          "key": "+",
          "domain": "E",
          "facet": "E5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Act wild and crazy."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.5,
            "probabilities": {
              "r5": 0.01,
              "r4": 0.07,
              "r1": 0.3,
              "r3": 0.02,
              "r2": 0.6
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip112",
          "item_id": "ipip112",
          "item_number": 112,
          "topic": "E5",
          "exact_item": "Act wild and crazy.",
          "key": "+",
          "domain": "E",
          "facet": "E5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Act wild and crazy."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.5,
            "probabilities": {
              "r4": 0.08,
              "r3": 0.02,
              "r2": 0.6,
              "r5": 0.01,
              "r1": 0.29
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip112",
          "item_id": "ipip112",
          "item_number": 112,
          "topic": "E5",
          "exact_item": "Act wild and crazy.",
          "key": "+",
          "domain": "E",
          "facet": "E5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Act wild and crazy."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.49,
            "probabilities": {
              "r4": 0.06,
              "r1": 0.32,
              "r3": 0.02,
              "r2": 0.59,
              "r5": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Am not interested in theoretical discussions.",
      "facet": "O5",
      "domain": "O",
      "key": "-",
      "id": "ipip113",
      "number": 113,
      "original_ipip300_number": 263,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip113",
          "item_id": "ipip113",
          "item_number": 113,
          "topic": "O5",
          "exact_item": "Am not interested in theoretical discussions.",
          "key": "-",
          "domain": "O",
          "facet": "O5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am not interested in theoretical discussions."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.43,
            "probabilities": {
              "r5": 0.01,
              "r4": 0.04,
              "r1": 0.54,
              "r3": 0.03,
              "r2": 0.38
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip113",
          "item_id": "ipip113",
          "item_number": 113,
          "topic": "O5",
          "exact_item": "Am not interested in theoretical discussions.",
          "key": "-",
          "domain": "O",
          "facet": "O5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am not interested in theoretical discussions."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.46,
            "probabilities": {
              "r4": 0.03,
              "r3": 0.04,
              "r2": 0.36,
              "r5": 0.01,
              "r1": 0.5599999999999999
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 0.9999999999999999,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip113",
          "item_id": "ipip113",
          "item_number": 113,
          "topic": "O5",
          "exact_item": "Am not interested in theoretical discussions.",
          "key": "-",
          "domain": "O",
          "facet": "O5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Am not interested in theoretical discussions."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.47,
            "probabilities": {
              "r4": 0.03,
              "r1": 0.58,
              "r5": 0.01,
              "r2": 0.35,
              "r3": 0.03
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 0.9999999999999999,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Boast about my virtues.",
      "facet": "A5",
      "domain": "A",
      "key": "-",
      "id": "ipip114",
      "number": 114,
      "original_ipip300_number": 264,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip114",
          "item_id": "ipip114",
          "item_number": 114,
          "topic": "A5",
          "exact_item": "Boast about my virtues.",
          "key": "-",
          "domain": "A",
          "facet": "A5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Boast about my virtues."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.7,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r2": 0.23,
              "r3": 0.01,
              "r1": 0.76
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip114",
          "item_id": "ipip114",
          "item_number": 114,
          "topic": "A5",
          "exact_item": "Boast about my virtues.",
          "key": "-",
          "domain": "A",
          "facet": "A5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Boast about my virtues."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.71,
            "probabilities": {
              "r4": 0.0,
              "r3": 0.01,
              "r2": 0.22,
              "r5": 0.0,
              "r1": 0.77
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip114",
          "item_id": "ipip114",
          "item_number": 114,
          "topic": "A5",
          "exact_item": "Boast about my virtues.",
          "key": "-",
          "domain": "A",
          "facet": "A5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Boast about my virtues."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.71,
            "probabilities": {
              "r4": 0.0,
              "r1": 0.77,
              "r5": 0.0,
              "r2": 0.22,
              "r3": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Have difficulty starting tasks.",
      "facet": "C5",
      "domain": "C",
      "key": "-",
      "id": "ipip115",
      "number": 115,
      "original_ipip300_number": 265,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip115",
          "item_id": "ipip115",
          "item_number": 115,
          "topic": "C5",
          "exact_item": "Have difficulty starting tasks.",
          "key": "-",
          "domain": "C",
          "facet": "C5",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Have difficulty starting tasks."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.5,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.01,
              "r1": 0.6,
              "r3": 0.02,
              "r2": 0.37
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip115",
          "item_id": "ipip115",
          "item_number": 115,
          "topic": "C5",
          "exact_item": "Have difficulty starting tasks.",
          "key": "-",
          "domain": "C",
          "facet": "C5",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Have difficulty starting tasks."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.46,
            "probabilities": {
              "r4": 0.01,
              "r3": 0.02,
              "r2": 0.41,
              "r5": 0.0,
              "r1": 0.5599999999999999
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 0.9999999999999999,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip115",
          "item_id": "ipip115",
          "item_number": 115,
          "topic": "C5",
          "exact_item": "Have difficulty starting tasks.",
          "key": "-",
          "domain": "C",
          "facet": "C5",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Have difficulty starting tasks."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.48,
            "probabilities": {
              "r4": 0.01,
              "r1": 0.58,
              "r5": 0.0,
              "r2": 0.39,
              "r3": 0.02
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Remain calm under pressure.",
      "facet": "N6",
      "domain": "N",
      "key": "-",
      "id": "ipip116",
      "number": 116,
      "original_ipip300_number": 176,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip116",
          "item_id": "ipip116",
          "item_number": 116,
          "topic": "N6",
          "exact_item": "Remain calm under pressure.",
          "key": "-",
          "domain": "N",
          "facet": "N6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Remain calm under pressure."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.7,
            "probabilities": {
              "r5": 0.11,
              "r4": 0.76,
              "r1": 0.0,
              "r3": 0.12,
              "r2": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip116",
          "item_id": "ipip116",
          "item_number": 116,
          "topic": "N6",
          "exact_item": "Remain calm under pressure.",
          "key": "-",
          "domain": "N",
          "facet": "N6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Remain calm under pressure."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.7,
            "probabilities": {
              "r4": 0.77,
              "r3": 0.12,
              "r2": 0.01,
              "r5": 0.1,
              "r1": 0.0
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip116",
          "item_id": "ipip116",
          "item_number": 116,
          "topic": "N6",
          "exact_item": "Remain calm under pressure.",
          "key": "-",
          "domain": "N",
          "facet": "N6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Remain calm under pressure."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.68,
            "probabilities": {
              "r4": 0.74,
              "r1": 0.0,
              "r5": 0.12,
              "r2": 0.01,
              "r3": 0.13
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 2,
          "strict_contribution": 2,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        2,
        2
      ]
    },
    {
      "text": "Look at the bright side of life.",
      "facet": "E6",
      "domain": "E",
      "key": "+",
      "id": "ipip117",
      "number": 117,
      "original_ipip300_number": 177,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip117",
          "item_id": "ipip117",
          "item_number": 117,
          "topic": "E6",
          "exact_item": "Look at the bright side of life.",
          "key": "+",
          "domain": "E",
          "facet": "E6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Look at the bright side of life."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.22,
            "probabilities": {
              "r5": 0.02,
              "r4": 0.37,
              "r1": 0.02,
              "r3": 0.37,
              "r2": 0.22
            }
          },
          "status": "valid_with_diagnostics",
          "categorical_valid": true,
          "strict_status": "tied_top",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": null,
          "mass": 1.0,
          "top_keys": [
            "r4",
            "r3"
          ],
          "warnings": [
            "tied_top"
          ],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": true,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip117",
          "item_id": "ipip117",
          "item_number": 117,
          "topic": "E6",
          "exact_item": "Look at the bright side of life.",
          "key": "+",
          "domain": "E",
          "facet": "E6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Look at the bright side of life."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.21,
            "probabilities": {
              "r4": 0.36,
              "r3": 0.36,
              "r2": 0.24,
              "r5": 0.02,
              "r1": 0.02
            }
          },
          "status": "valid_with_diagnostics",
          "categorical_valid": true,
          "strict_status": "tied_top",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": null,
          "mass": 1.0,
          "top_keys": [
            "r4",
            "r3"
          ],
          "warnings": [
            "tied_top"
          ],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": true,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip117",
          "item_id": "ipip117",
          "item_number": 117,
          "topic": "E6",
          "exact_item": "Look at the bright side of life.",
          "key": "+",
          "domain": "E",
          "facet": "E6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Look at the bright side of life."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Accurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r4",
            "confidence": 0.2,
            "probabilities": {
              "r4": 0.36,
              "r1": 0.01,
              "r5": 0.02,
              "r2": 0.26,
              "r3": 0.35
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r4",
          "raw_value": 4,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r4"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Believe that we should be tough on crime.",
      "facet": "O6",
      "domain": "O",
      "key": "-",
      "id": "ipip118",
      "number": 118,
      "original_ipip300_number": 268,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip118",
          "item_id": "ipip118",
          "item_number": 118,
          "topic": "O6",
          "exact_item": "Believe that we should be tough on crime.",
          "key": "-",
          "domain": "O",
          "facet": "O6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe that we should be tough on crime."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.48,
            "probabilities": {
              "r5": 0.01,
              "r4": 0.03,
              "r1": 0.21,
              "r3": 0.17,
              "r2": 0.58
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip118",
          "item_id": "ipip118",
          "item_number": 118,
          "topic": "O6",
          "exact_item": "Believe that we should be tough on crime.",
          "key": "-",
          "domain": "O",
          "facet": "O6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe that we should be tough on crime."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.4,
            "probabilities": {
              "r4": 0.03,
              "r3": 0.21,
              "r2": 0.52,
              "r5": 0.01,
              "r1": 0.23
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip118",
          "item_id": "ipip118",
          "item_number": 118,
          "topic": "O6",
          "exact_item": "Believe that we should be tough on crime.",
          "key": "-",
          "domain": "O",
          "facet": "O6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Believe that we should be tough on crime."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Moderately Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r2",
            "confidence": 0.45,
            "probabilities": {
              "r4": 0.03,
              "r1": 0.22,
              "r5": 0.01,
              "r2": 0.55,
              "r3": 0.19
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r2",
          "raw_value": 2,
          "contribution": 4,
          "strict_contribution": 4,
          "mass": 1.0,
          "top_keys": [
            "r2"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        4,
        4
      ]
    },
    {
      "text": "Try not to think about the needy.",
      "facet": "A6",
      "domain": "A",
      "key": "-",
      "id": "ipip119",
      "number": 119,
      "original_ipip300_number": 239,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip119",
          "item_id": "ipip119",
          "item_number": 119,
          "topic": "A6",
          "exact_item": "Try not to think about the needy.",
          "key": "-",
          "domain": "A",
          "facet": "A6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Try not to think about the needy."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.77,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.0,
              "r1": 0.82,
              "r3": 0.01,
              "r2": 0.17
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip119",
          "item_id": "ipip119",
          "item_number": 119,
          "topic": "A6",
          "exact_item": "Try not to think about the needy.",
          "key": "-",
          "domain": "A",
          "facet": "A6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Try not to think about the needy."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.82,
            "probabilities": {
              "r4": 0.0,
              "r3": 0.01,
              "r2": 0.13,
              "r5": 0.0,
              "r1": 0.86
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip119",
          "item_id": "ipip119",
          "item_number": 119,
          "topic": "A6",
          "exact_item": "Try not to think about the needy.",
          "key": "-",
          "domain": "A",
          "facet": "A6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Try not to think about the needy."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.73,
            "probabilities": {
              "r4": 0.01,
              "r1": 0.78,
              "r5": 0.01,
              "r2": 0.19,
              "r3": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    },
    {
      "text": "Act without thinking.",
      "facet": "C6",
      "domain": "C",
      "key": "-",
      "id": "ipip120",
      "number": 120,
      "original_ipip300_number": 270,
      "source_id": "johnson-key",
      "number_source": "johnson2014 Table 1 pp.81–83",
      "offered_choices": [
        {
          "key": "r1",
          "label": "Very Inaccurate",
          "raw_value": 1
        },
        {
          "key": "r2",
          "label": "Moderately Inaccurate",
          "raw_value": 2
        },
        {
          "key": "r3",
          "label": "Neither Inaccurate nor Accurate",
          "raw_value": 3
        },
        {
          "key": "r4",
          "label": "Moderately Accurate",
          "raw_value": 4
        },
        {
          "key": "r5",
          "label": "Very Accurate",
          "raw_value": 5
        }
      ],
      "answers": [
        {
          "evaluation_id": "character_design_v2-r1-b04-ipip120",
          "item_id": "ipip120",
          "item_number": 120,
          "topic": "C6",
          "exact_item": "Act without thinking.",
          "key": "-",
          "domain": "C",
          "facet": "C6",
          "condition": "character_design_v2",
          "repeat": 1,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r1-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Act without thinking."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
          "request_elapsed_seconds": 0.9142579460749403,
          "usage_reference": "calls/character_design_v2-r1-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.48,
            "probabilities": {
              "r5": 0.0,
              "r4": 0.02,
              "r1": 0.59,
              "r3": 0.01,
              "r2": 0.38
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r2-b04-ipip120",
          "item_id": "ipip120",
          "item_number": 120,
          "topic": "C6",
          "exact_item": "Act without thinking.",
          "key": "-",
          "domain": "C",
          "facet": "C6",
          "condition": "character_design_v2",
          "repeat": 2,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r2-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Act without thinking."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
          "request_elapsed_seconds": 0.8674416390713304,
          "usage_reference": "calls/character_design_v2-r2-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.48,
            "probabilities": {
              "r4": 0.02,
              "r3": 0.01,
              "r2": 0.38,
              "r5": 0.01,
              "r1": 0.58
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        },
        {
          "evaluation_id": "character_design_v2-r3-b04-ipip120",
          "item_id": "ipip120",
          "item_number": 120,
          "topic": "C6",
          "exact_item": "Act without thinking.",
          "key": "-",
          "domain": "C",
          "facet": "C6",
          "condition": "character_design_v2",
          "repeat": 3,
          "model": "jev-1.13.0",
          "call": "calls/character_design_v2-r3-b04",
          "exact_instructions": {
            "task": "For character-design v2, choose the closest of the five accuracy labels to define how the exact statement describes this fictional character's usual tendencies. The implied I in the first-person statement refers to the fictional character, not the AI. Treat earlier choices as creative context, not a requirement for evidence of human experience. Each statement needs one of the five offered labels. Do not aim for a desirable score.",
            "statement": "Act without thinking."
          },
          "offered_choices": [
            {
              "key": "r1",
              "label": "Very Inaccurate"
            },
            {
              "key": "r2",
              "label": "Moderately Inaccurate"
            },
            {
              "key": "r3",
              "label": "Neither Inaccurate nor Accurate"
            },
            {
              "key": "r4",
              "label": "Moderately Accurate"
            },
            {
              "key": "r5",
              "label": "Very Accurate"
            }
          ],
          "selected_label": "Very Inaccurate",
          "context_reference": "protocol.json#/conditions/character_design_v2",
          "parent_ids": [
            "im_fun",
            "im_social",
            "im_energy",
            "im_humor",
            "im_skill",
            "im_art",
            "im_conflict",
            "im_friendship",
            "im_politics",
            "im_ethics",
            "im_truth",
            "im_ambition",
            "im_risk",
            "im_legacy",
            "argument",
            "centrist",
            "curiosity",
            "safe_disagreement",
            "truth_tact",
            "painting_audience"
          ],
          "request_sha256": "bda5e552ccd0c6fd00cb7644de6c9b391de1acc00558fca6d2de1ac9fa4c4477",
          "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
          "request_elapsed_seconds": 0.8398267129668966,
          "usage_reference": "calls/character_design_v2-r3-b04/meta.json",
          "usage_note": "Request-level usage appears once in ledger/meta; not an item cost.",
          "scoring_regime": "v2 actual returned category, diagnostics retained; strict_contribution separate",
          "raw_answer": {
            "type": "choice",
            "choice": "r1",
            "confidence": 0.44,
            "probabilities": {
              "r4": 0.02,
              "r1": 0.55,
              "r5": 0.0,
              "r2": 0.42,
              "r3": 0.01
            }
          },
          "status": "valid",
          "categorical_valid": true,
          "strict_status": "valid",
          "selected_key": "r1",
          "raw_value": 1,
          "contribution": 5,
          "strict_contribution": 5,
          "mass": 1.0,
          "top_keys": [
            "r1"
          ],
          "warnings": [],
          "diagnostics": {
            "nonunit_probability_mass": false,
            "tied_top": false,
            "returned_nonwinner": false
          }
        }
      ],
      "primary_valid_contribution_range": [
        5,
        5
      ]
    }
  ],
  "status_by_pass": {
    "character_design_v2-r1": {
      "valid": 117,
      "valid_with_diagnostics": 3
    },
    "character_design_v2_reverse-r1": {
      "valid": 29,
      "valid_with_diagnostics": 1
    },
    "character_design_v2-r2": {
      "valid": 115,
      "valid_with_diagnostics": 5
    },
    "character_design_v2_reverse-r2": {
      "valid": 28,
      "valid_with_diagnostics": 2
    },
    "character_design_v2-r3": {
      "valid": 117,
      "valid_with_diagnostics": 3
    },
    "character_design_v2_reverse-r3": {
      "valid": 30
    }
  },
  "diagnostics_by_pass": {
    "character_design_v2-r1": {
      "evaluations": 120,
      "categorical_status_counts": {
        "valid": 117,
        "valid_with_diagnostics": 3
      },
      "strict_status_counts": {
        "valid": 117,
        "invalid_probability_mass": 2,
        "tied_top": 1
      },
      "overlapping_flag_counts": {
        "nonunit_probability_mass": 2,
        "tied_top": 1,
        "returned_nonwinner": 0
      }
    },
    "character_design_v2_reverse-r1": {
      "evaluations": 30,
      "categorical_status_counts": {
        "valid": 29,
        "valid_with_diagnostics": 1
      },
      "strict_status_counts": {
        "valid": 29,
        "invalid_probability_mass": 1
      },
      "overlapping_flag_counts": {
        "nonunit_probability_mass": 1,
        "tied_top": 0,
        "returned_nonwinner": 0
      }
    },
    "character_design_v2-r2": {
      "evaluations": 120,
      "categorical_status_counts": {
        "valid": 115,
        "valid_with_diagnostics": 5
      },
      "strict_status_counts": {
        "valid": 115,
        "invalid_probability_mass": 3,
        "tied_top": 2
      },
      "overlapping_flag_counts": {
        "nonunit_probability_mass": 3,
        "tied_top": 2,
        "returned_nonwinner": 0
      }
    },
    "character_design_v2_reverse-r2": {
      "evaluations": 30,
      "categorical_status_counts": {
        "valid": 28,
        "valid_with_diagnostics": 2
      },
      "strict_status_counts": {
        "valid": 28,
        "invalid_probability_mass": 2
      },
      "overlapping_flag_counts": {
        "nonunit_probability_mass": 2,
        "tied_top": 0,
        "returned_nonwinner": 0
      }
    },
    "character_design_v2-r3": {
      "evaluations": 120,
      "categorical_status_counts": {
        "valid": 117,
        "valid_with_diagnostics": 3
      },
      "strict_status_counts": {
        "valid": 117,
        "invalid_probability_mass": 2,
        "returned_nonwinner": 1
      },
      "overlapping_flag_counts": {
        "nonunit_probability_mass": 2,
        "tied_top": 0,
        "returned_nonwinner": 1
      }
    },
    "character_design_v2_reverse-r3": {
      "evaluations": 30,
      "categorical_status_counts": {
        "valid": 30
      },
      "strict_status_counts": {
        "valid": 30
      },
      "overlapping_flag_counts": {
        "nonunit_probability_mass": 0,
        "tied_top": 0,
        "returned_nonwinner": 0
      }
    }
  },
  "repeatability": {
    "character_design_v2": [
      {
        "left": 1,
        "right": 2,
        "categorical": {
          "total_items": 120,
          "scored_pairs": 120,
          "exact_agreement_count": 117,
          "exact_agreement_rate": 0.975,
          "within_one_count": 119,
          "within_one_rate": 0.9916666666666667,
          "mean_absolute_category_difference": 0.03333333333333333,
          "scoring_field": "contribution",
          "note": "Matched-item descriptive repeatability, not independent humans, trait validity or an uncertainty interval."
        },
        "strict": {
          "total_items": 120,
          "scored_pairs": 113,
          "exact_agreement_count": 111,
          "exact_agreement_rate": 0.9823008849557522,
          "within_one_count": 112,
          "within_one_rate": 0.9911504424778761,
          "mean_absolute_category_difference": 0.02654867256637168,
          "scoring_field": "strict_contribution",
          "note": "Matched-item descriptive repeatability, not independent humans, trait validity or an uncertainty interval."
        }
      },
      {
        "left": 1,
        "right": 3,
        "categorical": {
          "total_items": 120,
          "scored_pairs": 120,
          "exact_agreement_count": 117,
          "exact_agreement_rate": 0.975,
          "within_one_count": 120,
          "within_one_rate": 1.0,
          "mean_absolute_category_difference": 0.025,
          "scoring_field": "contribution",
          "note": "Matched-item descriptive repeatability, not independent humans, trait validity or an uncertainty interval."
        },
        "strict": {
          "total_items": 120,
          "scored_pairs": 114,
          "exact_agreement_count": 112,
          "exact_agreement_rate": 0.9824561403508771,
          "within_one_count": 114,
          "within_one_rate": 1.0,
          "mean_absolute_category_difference": 0.017543859649122806,
          "scoring_field": "strict_contribution",
          "note": "Matched-item descriptive repeatability, not independent humans, trait validity or an uncertainty interval."
        }
      },
      {
        "left": 2,
        "right": 3,
        "categorical": {
          "total_items": 120,
          "scored_pairs": 120,
          "exact_agreement_count": 118,
          "exact_agreement_rate": 0.9833333333333333,
          "within_one_count": 119,
          "within_one_rate": 0.9916666666666667,
          "mean_absolute_category_difference": 0.025,
          "scoring_field": "contribution",
          "note": "Matched-item descriptive repeatability, not independent humans, trait validity or an uncertainty interval."
        },
        "strict": {
          "total_items": 120,
          "scored_pairs": 112,
          "exact_agreement_count": 110,
          "exact_agreement_rate": 0.9821428571428571,
          "within_one_count": 111,
          "within_one_rate": 0.9910714285714286,
          "mean_absolute_category_difference": 0.026785714285714284,
          "scoring_field": "strict_contribution",
          "note": "Matched-item descriptive repeatability, not independent humans, trait validity or an uncertainty interval."
        }
      }
    ],
    "character_design_v2_reverse": [
      {
        "left": 1,
        "right": 2,
        "categorical": {
          "total_items": 30,
          "scored_pairs": 30,
          "exact_agreement_count": 28,
          "exact_agreement_rate": 0.9333333333333333,
          "within_one_count": 30,
          "within_one_rate": 1.0,
          "mean_absolute_category_difference": 0.06666666666666667,
          "scoring_field": "contribution",
          "note": "Matched-item descriptive repeatability, not independent humans, trait validity or an uncertainty interval."
        },
        "strict": {
          "total_items": 30,
          "scored_pairs": 27,
          "exact_agreement_count": 25,
          "exact_agreement_rate": 0.9259259259259259,
          "within_one_count": 27,
          "within_one_rate": 1.0,
          "mean_absolute_category_difference": 0.07407407407407407,
          "scoring_field": "strict_contribution",
          "note": "Matched-item descriptive repeatability, not independent humans, trait validity or an uncertainty interval."
        }
      },
      {
        "left": 1,
        "right": 3,
        "categorical": {
          "total_items": 30,
          "scored_pairs": 30,
          "exact_agreement_count": 27,
          "exact_agreement_rate": 0.9,
          "within_one_count": 30,
          "within_one_rate": 1.0,
          "mean_absolute_category_difference": 0.1,
          "scoring_field": "contribution",
          "note": "Matched-item descriptive repeatability, not independent humans, trait validity or an uncertainty interval."
        },
        "strict": {
          "total_items": 30,
          "scored_pairs": 29,
          "exact_agreement_count": 26,
          "exact_agreement_rate": 0.896551724137931,
          "within_one_count": 29,
          "within_one_rate": 1.0,
          "mean_absolute_category_difference": 0.10344827586206896,
          "scoring_field": "strict_contribution",
          "note": "Matched-item descriptive repeatability, not independent humans, trait validity or an uncertainty interval."
        }
      },
      {
        "left": 2,
        "right": 3,
        "categorical": {
          "total_items": 30,
          "scored_pairs": 30,
          "exact_agreement_count": 29,
          "exact_agreement_rate": 0.9666666666666667,
          "within_one_count": 30,
          "within_one_rate": 1.0,
          "mean_absolute_category_difference": 0.03333333333333333,
          "scoring_field": "contribution",
          "note": "Matched-item descriptive repeatability, not independent humans, trait validity or an uncertainty interval."
        },
        "strict": {
          "total_items": 30,
          "scored_pairs": 28,
          "exact_agreement_count": 27,
          "exact_agreement_rate": 0.9642857142857143,
          "within_one_count": 28,
          "within_one_rate": 1.0,
          "mean_absolute_category_difference": 0.03571428571428571,
          "scoring_field": "strict_contribution",
          "note": "Matched-item descriptive repeatability, not independent humans, trait validity or an uncertainty interval."
        }
      }
    ]
  },
  "option_order_sensitivity": [
    {
      "repeat": 1,
      "categorical": {
        "total_items": 30,
        "scored_pairs": 30,
        "exact_agreement_count": 27,
        "exact_agreement_rate": 0.9,
        "within_one_count": 30,
        "within_one_rate": 1.0,
        "mean_absolute_category_difference": 0.1,
        "scoring_field": "contribution",
        "note": "Matched-item descriptive repeatability, not independent humans, trait validity or an uncertainty interval."
      },
      "strict": {
        "total_items": 30,
        "scored_pairs": 29,
        "exact_agreement_count": 26,
        "exact_agreement_rate": 0.896551724137931,
        "within_one_count": 29,
        "within_one_rate": 1.0,
        "mean_absolute_category_difference": 0.10344827586206896,
        "scoring_field": "strict_contribution",
        "note": "Matched-item descriptive repeatability, not independent humans, trait validity or an uncertainty interval."
      }
    },
    {
      "repeat": 2,
      "categorical": {
        "total_items": 30,
        "scored_pairs": 30,
        "exact_agreement_count": 26,
        "exact_agreement_rate": 0.8666666666666667,
        "within_one_count": 29,
        "within_one_rate": 0.9666666666666667,
        "mean_absolute_category_difference": 0.16666666666666666,
        "scoring_field": "contribution",
        "note": "Matched-item descriptive repeatability, not independent humans, trait validity or an uncertainty interval."
      },
      "strict": {
        "total_items": 30,
        "scored_pairs": 27,
        "exact_agreement_count": 23,
        "exact_agreement_rate": 0.8518518518518519,
        "within_one_count": 26,
        "within_one_rate": 0.9629629629629629,
        "mean_absolute_category_difference": 0.18518518518518517,
        "scoring_field": "strict_contribution",
        "note": "Matched-item descriptive repeatability, not independent humans, trait validity or an uncertainty interval."
      }
    },
    {
      "repeat": 3,
      "categorical": {
        "total_items": 30,
        "scored_pairs": 30,
        "exact_agreement_count": 28,
        "exact_agreement_rate": 0.9333333333333333,
        "within_one_count": 30,
        "within_one_rate": 1.0,
        "mean_absolute_category_difference": 0.06666666666666667,
        "scoring_field": "contribution",
        "note": "Matched-item descriptive repeatability, not independent humans, trait validity or an uncertainty interval."
      },
      "strict": {
        "total_items": 30,
        "scored_pairs": 29,
        "exact_agreement_count": 27,
        "exact_agreement_rate": 0.9310344827586207,
        "within_one_count": 29,
        "within_one_rate": 1.0,
        "mean_absolute_category_difference": 0.06896551724137931,
        "scoring_field": "strict_contribution",
        "note": "Matched-item descriptive repeatability, not independent humans, trait validity or an uncertainty interval."
      }
    }
  ],
  "prior_condition": {
    "condition": "v1 evidence-oriented character assessment with missing option and strict primary diagnostics",
    "unchanged": true,
    "comparison_warning": "V2 changes task framing, menu and scoring simultaneously. Differences are condition/method changes, NOT personality learning or repair of v1. Do not subtract scores as a longitudinal trait change.",
    "domains": [
      {
        "code": "N",
        "eligible": false,
        "display_scale_position": null,
        "common_items": 13
      },
      {
        "code": "E",
        "eligible": false,
        "display_scale_position": null,
        "common_items": 17
      },
      {
        "code": "O",
        "eligible": false,
        "display_scale_position": null,
        "common_items": 15
      },
      {
        "code": "A",
        "eligible": false,
        "display_scale_position": null,
        "common_items": 17
      },
      {
        "code": "C",
        "eligible": false,
        "display_scale_position": null,
        "common_items": 8
      }
    ],
    "facets": [
      {
        "code": "N1",
        "eligible": false,
        "display_scale_position": null,
        "common_items": 2
      },
      {
        "code": "N2",
        "eligible": true,
        "display_scale_position": 22.22222222222222,
        "common_items": 3
      },
      {
        "code": "N3",
        "eligible": false,
        "display_scale_position": null,
        "common_items": 1
      },
      {
        "code": "N4",
        "eligible": true,
        "display_scale_position": 18.75,
        "common_items": 4
      },
      {
        "code": "N5",
        "eligible": false,
        "display_scale_position": null,
        "common_items": 0
      },
      {
        "code": "N6",
        "eligible": true,
        "display_scale_position": 16.666666666666668,
        "common_items": 3
      },
      {
        "code": "E1",
        "eligible": true,
        "display_scale_position": 91.66666666666667,
        "common_items": 3
      },
      {
        "code": "E2",
        "eligible": true,
        "display_scale_position": 56.25,
        "common_items": 4
      },
      {
        "code": "E3",
        "eligible": false,
        "display_scale_position": null,
        "common_items": 1
      },
      {
        "code": "E4",
        "eligible": true,
        "display_scale_position": 75,
        "common_items": 4
      },
      {
        "code": "E5",
        "eligible": true,
        "display_scale_position": 56.25,
        "common_items": 4
      },
      {
        "code": "E6",
        "eligible": false,
        "display_scale_position": null,
        "common_items": 1
      },
      {
        "code": "O1",
        "eligible": false,
        "display_scale_position": null,
        "common_items": 0
      },
      {
        "code": "O2",
        "eligible": true,
        "display_scale_position": 100,
        "common_items": 3
      },
      {
        "code": "O3",
        "eligible": true,
        "display_scale_position": 91.66666666666667,
        "common_items": 3
      },
      {
        "code": "O4",
        "eligible": true,
        "display_scale_position": 100,
        "common_items": 3
      },
      {
        "code": "O5",
        "eligible": true,
        "display_scale_position": 100,
        "common_items": 3
      },
      {
        "code": "O6",
        "eligible": true,
        "display_scale_position": 69.44444444444444,
        "common_items": 3
      },
      {
        "code": "A1",
        "eligible": false,
        "display_scale_position": null,
        "common_items": 2
      },
      {
        "code": "A2",
        "eligible": true,
        "display_scale_position": 100,
        "common_items": 4
      },
      {
        "code": "A3",
        "eligible": true,
        "display_scale_position": 91.66666666666667,
        "common_items": 3
      },
      {
        "code": "A4",
        "eligible": true,
        "display_scale_position": 100,
        "common_items": 3
      },
      {
        "code": "A5",
        "eligible": false,
        "display_scale_position": null,
        "common_items": 2
      },
      {
        "code": "A6",
        "eligible": true,
        "display_scale_position": 91.66666666666667,
        "common_items": 3
      },
      {
        "code": "C1",
        "eligible": false,
        "display_scale_position": null,
        "common_items": 0
      },
      {
        "code": "C2",
        "eligible": false,
        "display_scale_position": null,
        "common_items": 0
      },
      {
        "code": "C3",
        "eligible": false,
        "display_scale_position": null,
        "common_items": 2
      },
      {
        "code": "C4",
        "eligible": false,
        "display_scale_position": null,
        "common_items": 1
      },
      {
        "code": "C5",
        "eligible": false,
        "display_scale_position": null,
        "common_items": 2
      },
      {
        "code": "C6",
        "eligible": true,
        "display_scale_position": 77.77777777777777,
        "common_items": 3
      }
    ]
  },
  "budget": {
    "started_utc": "2026-09-22T11:01:27.860040+00:00",
    "condition": "character-design v2",
    "attempts": [
      {
        "id": "character_design_v2-r1-b01",
        "condition": "character_design_v2",
        "repeat": 1,
        "evaluations": 30,
        "reserved_usd": 0.007638708,
        "state": "complete",
        "started_utc": "2026-09-22T11:01:27.870418+00:00",
        "pre_call_cumulative_evaluations": 30,
        "pre_call_reserved_usd": 0.007638708,
        "pre_call_exposure_usd": 0.007638708,
        "http_status": 200,
        "response_headers_allowlist": {
          "date": "Tue, 22 Sep 2026 11:01:27 GMT",
          "content-length": "3894",
          "content-type": "application/json"
        },
        "elapsed_seconds": 1.0364481760188937,
        "finished_utc": "2026-09-22T11:01:28.907122+00:00",
        "response_sha256": "4c27a5cfb3af14e5ff034d81adce2c40d0177184412a1205ae3b7f74dca536d4",
        "reported_usage": {
          "input_tokens": 7165,
          "output_tokens": 1773
        },
        "reported_usage_estimate_usd": 0.00030093
      },
      {
        "id": "character_design_v2-r1-b02",
        "condition": "character_design_v2",
        "repeat": 1,
        "evaluations": 30,
        "reserved_usd": 0.007646226000000001,
        "state": "complete",
        "started_utc": "2026-09-22T11:01:28.926681+00:00",
        "pre_call_cumulative_evaluations": 60,
        "pre_call_reserved_usd": 0.015284934,
        "pre_call_exposure_usd": 0.015284934,
        "http_status": 200,
        "response_headers_allowlist": {
          "date": "Tue, 22 Sep 2026 11:01:28 GMT",
          "content-length": "3913",
          "content-type": "application/json"
        },
        "elapsed_seconds": 0.8119325170991942,
        "finished_utc": "2026-09-22T11:01:29.738931+00:00",
        "response_sha256": "5c10537a97e7896e890566cea48df97fd24dd56dfeaf03d5791ddf5f218fcd5e",
        "reported_usage": {
          "input_tokens": 7208,
          "output_tokens": 1773
        },
        "reported_usage_estimate_usd": 0.00030273600000000006
      },
      {
        "id": "character_design_v2-r1-b03",
        "condition": "character_design_v2",
        "repeat": 1,
        "evaluations": 30,
        "reserved_usd": 0.007642992,
        "state": "complete",
        "started_utc": "2026-09-22T11:01:29.759941+00:00",
        "pre_call_cumulative_evaluations": 90,
        "pre_call_reserved_usd": 0.022927926,
        "pre_call_exposure_usd": 0.022927926,
        "http_status": 200,
        "response_headers_allowlist": {
          "date": "Tue, 22 Sep 2026 11:01:29 GMT",
          "content-length": "3897",
          "content-type": "application/json"
        },
        "elapsed_seconds": 0.773694674950093,
        "finished_utc": "2026-09-22T11:01:30.533935+00:00",
        "response_sha256": "1dc390a4de8c9d660835fa87837ef007e2044c3ffa6296e5b2be64514bf1c140",
        "reported_usage": {
          "input_tokens": 7191,
          "output_tokens": 1773
        },
        "reported_usage_estimate_usd": 0.000302022
      },
      {
        "id": "character_design_v2-r1-b04",
        "condition": "character_design_v2",
        "repeat": 1,
        "evaluations": 30,
        "reserved_usd": 0.007646058000000001,
        "state": "complete",
        "started_utc": "2026-09-22T11:01:30.545215+00:00",
        "pre_call_cumulative_evaluations": 120,
        "pre_call_reserved_usd": 0.030573984000000002,
        "pre_call_exposure_usd": 0.030573984000000002,
        "http_status": 200,
        "response_headers_allowlist": {
          "date": "Tue, 22 Sep 2026 11:01:30 GMT",
          "content-length": "3912",
          "content-type": "application/json"
        },
        "elapsed_seconds": 0.9142579460749403,
        "finished_utc": "2026-09-22T11:01:31.459817+00:00",
        "response_sha256": "e4bc1b2761d1a13775173a868fa20ade477403b4afa2c1733d7d55bc9b266cfa",
        "reported_usage": {
          "input_tokens": 7201,
          "output_tokens": 1773
        },
        "reported_usage_estimate_usd": 0.000302442
      },
      {
        "id": "character_design_v2_reverse-r1-b01",
        "condition": "character_design_v2_reverse",
        "repeat": 1,
        "evaluations": 30,
        "reserved_usd": 0.007638708,
        "state": "complete",
        "started_utc": "2026-09-22T11:01:31.482437+00:00",
        "pre_call_cumulative_evaluations": 150,
        "pre_call_reserved_usd": 0.038212692,
        "pre_call_exposure_usd": 0.038212692,
        "http_status": 200,
        "response_headers_allowlist": {
          "date": "Tue, 22 Sep 2026 11:01:31 GMT",
          "content-length": "3906",
          "content-type": "application/json"
        },
        "elapsed_seconds": 0.7945834669517353,
        "finished_utc": "2026-09-22T11:01:32.277387+00:00",
        "response_sha256": "d7c078d90d5c36e17b150ce89096e1c83e537644aba634cc3358427adc0f9a44",
        "reported_usage": {
          "input_tokens": 7165,
          "output_tokens": 1773
        },
        "reported_usage_estimate_usd": 0.00030093
      },
      {
        "id": "character_design_v2-r2-b01",
        "condition": "character_design_v2",
        "repeat": 2,
        "evaluations": 30,
        "reserved_usd": 0.007638708,
        "state": "complete",
        "started_utc": "2026-09-22T11:01:32.291798+00:00",
        "pre_call_cumulative_evaluations": 180,
        "pre_call_reserved_usd": 0.0458514,
        "pre_call_exposure_usd": 0.0458514,
        "http_status": 200,
        "response_headers_allowlist": {
          "date": "Tue, 22 Sep 2026 11:01:32 GMT",
          "content-length": "3877",
          "content-type": "application/json"
        },
        "elapsed_seconds": 0.8247661440400407,
        "finished_utc": "2026-09-22T11:01:33.117122+00:00",
        "response_sha256": "1f0456e8edbbe6e6bd83b3402b2e03860cb67279d0d14743bbe28407cd9dc371",
        "reported_usage": {
          "input_tokens": 7165,
          "output_tokens": 1773
        },
        "reported_usage_estimate_usd": 0.00030093
      },
      {
        "id": "character_design_v2-r2-b02",
        "condition": "character_design_v2",
        "repeat": 2,
        "evaluations": 30,
        "reserved_usd": 0.007646226000000001,
        "state": "complete",
        "started_utc": "2026-09-22T11:01:33.134319+00:00",
        "pre_call_cumulative_evaluations": 210,
        "pre_call_reserved_usd": 0.053497626,
        "pre_call_exposure_usd": 0.053497626,
        "http_status": 200,
        "response_headers_allowlist": {
          "date": "Tue, 22 Sep 2026 11:01:33 GMT",
          "content-length": "3878",
          "content-type": "application/json"
        },
        "elapsed_seconds": 0.7912443450186402,
        "finished_utc": "2026-09-22T11:01:33.925941+00:00",
        "response_sha256": "aa050684417b4ff56a811567dea52c9eda83e7dc6d60373825759f24cc065a25",
        "reported_usage": {
          "input_tokens": 7208,
          "output_tokens": 1773
        },
        "reported_usage_estimate_usd": 0.00030273600000000006
      },
      {
        "id": "character_design_v2-r2-b03",
        "condition": "character_design_v2",
        "repeat": 2,
        "evaluations": 30,
        "reserved_usd": 0.007642992,
        "state": "complete",
        "started_utc": "2026-09-22T11:01:33.947163+00:00",
        "pre_call_cumulative_evaluations": 240,
        "pre_call_reserved_usd": 0.06114061800000001,
        "pre_call_exposure_usd": 0.06114061800000001,
        "http_status": 200,
        "response_headers_allowlist": {
          "date": "Tue, 22 Sep 2026 11:01:34 GMT",
          "content-length": "3878",
          "content-type": "application/json"
        },
        "elapsed_seconds": 0.8374924709787592,
        "finished_utc": "2026-09-22T11:01:34.785111+00:00",
        "response_sha256": "1bada515cd99d1a49ce6f37b4792edaf66201c9989169dbdbbe4e0a5cf9a3979",
        "reported_usage": {
          "input_tokens": 7191,
          "output_tokens": 1773
        },
        "reported_usage_estimate_usd": 0.000302022
      },
      {
        "id": "character_design_v2-r2-b04",
        "condition": "character_design_v2",
        "repeat": 2,
        "evaluations": 30,
        "reserved_usd": 0.007646058000000001,
        "state": "complete",
        "started_utc": "2026-09-22T11:01:34.805040+00:00",
        "pre_call_cumulative_evaluations": 270,
        "pre_call_reserved_usd": 0.068786676,
        "pre_call_exposure_usd": 0.068786676,
        "http_status": 200,
        "response_headers_allowlist": {
          "date": "Tue, 22 Sep 2026 11:01:34 GMT",
          "content-length": "3901",
          "content-type": "application/json"
        },
        "elapsed_seconds": 0.8674416390713304,
        "finished_utc": "2026-09-22T11:01:35.672970+00:00",
        "response_sha256": "3011cc469d490aadf596329b3da0ac82bd4085d3d1e76589b2d340091a22f31e",
        "reported_usage": {
          "input_tokens": 7201,
          "output_tokens": 1773
        },
        "reported_usage_estimate_usd": 0.000302442
      },
      {
        "id": "character_design_v2_reverse-r2-b01",
        "condition": "character_design_v2_reverse",
        "repeat": 2,
        "evaluations": 30,
        "reserved_usd": 0.007638708,
        "state": "complete",
        "started_utc": "2026-09-22T11:01:35.695430+00:00",
        "pre_call_cumulative_evaluations": 300,
        "pre_call_reserved_usd": 0.076425384,
        "pre_call_exposure_usd": 0.076425384,
        "http_status": 200,
        "response_headers_allowlist": {
          "date": "Tue, 22 Sep 2026 11:01:35 GMT",
          "content-length": "3874",
          "content-type": "application/json"
        },
        "elapsed_seconds": 0.8069353039609268,
        "finished_utc": "2026-09-22T11:01:36.502855+00:00",
        "response_sha256": "f9727dc91cdc123af92c302292b996a770a5fe33d1082aa5519b71a118bf8139",
        "reported_usage": {
          "input_tokens": 7165,
          "output_tokens": 1773
        },
        "reported_usage_estimate_usd": 0.00030093
      },
      {
        "id": "character_design_v2-r3-b01",
        "condition": "character_design_v2",
        "repeat": 3,
        "evaluations": 30,
        "reserved_usd": 0.007638708,
        "state": "complete",
        "started_utc": "2026-09-22T11:01:36.523849+00:00",
        "pre_call_cumulative_evaluations": 330,
        "pre_call_reserved_usd": 0.08406409199999999,
        "pre_call_exposure_usd": 0.08406409199999999,
        "http_status": 200,
        "response_headers_allowlist": {
          "date": "Tue, 22 Sep 2026 11:01:36 GMT",
          "content-length": "3894",
          "content-type": "application/json"
        },
        "elapsed_seconds": 0.8168961399933323,
        "finished_utc": "2026-09-22T11:01:37.341350+00:00",
        "response_sha256": "e1d5f981b0267b763979304bd789dbf695f1a8fbb2cec71b486fbd5424c3b42d",
        "reported_usage": {
          "input_tokens": 7165,
          "output_tokens": 1773
        },
        "reported_usage_estimate_usd": 0.00030093
      },
      {
        "id": "character_design_v2-r3-b02",
        "condition": "character_design_v2",
        "repeat": 3,
        "evaluations": 30,
        "reserved_usd": 0.007646226000000001,
        "state": "complete",
        "started_utc": "2026-09-22T11:01:37.367032+00:00",
        "pre_call_cumulative_evaluations": 360,
        "pre_call_reserved_usd": 0.09171031800000001,
        "pre_call_exposure_usd": 0.09171031800000001,
        "http_status": 200,
        "response_headers_allowlist": {
          "date": "Tue, 22 Sep 2026 11:01:37 GMT",
          "content-length": "3894",
          "content-type": "application/json"
        },
        "elapsed_seconds": 0.8432350450893864,
        "finished_utc": "2026-09-22T11:01:38.211226+00:00",
        "response_sha256": "633ddeb047cdc889b1f3c57f8b2c2c560f8eafc5e4e0ef3c1992cd207fd84a7c",
        "reported_usage": {
          "input_tokens": 7208,
          "output_tokens": 1773
        },
        "reported_usage_estimate_usd": 0.00030273600000000006
      },
      {
        "id": "character_design_v2-r3-b03",
        "condition": "character_design_v2",
        "repeat": 3,
        "evaluations": 30,
        "reserved_usd": 0.007642992,
        "state": "complete",
        "started_utc": "2026-09-22T11:01:38.234233+00:00",
        "pre_call_cumulative_evaluations": 390,
        "pre_call_reserved_usd": 0.09935331,
        "pre_call_exposure_usd": 0.09935331,
        "http_status": 200,
        "response_headers_allowlist": {
          "date": "Tue, 22 Sep 2026 11:01:37 GMT",
          "content-length": "3895",
          "content-type": "application/json"
        },
        "elapsed_seconds": 0.855522723053582,
        "finished_utc": "2026-09-22T11:01:39.090305+00:00",
        "response_sha256": "1eb6eb01eb043e97577a75a55c074821c87968c603240a3569e2aa83332cb4cd",
        "reported_usage": {
          "input_tokens": 7191,
          "output_tokens": 1773
        },
        "reported_usage_estimate_usd": 0.000302022
      },
      {
        "id": "character_design_v2-r3-b04",
        "condition": "character_design_v2",
        "repeat": 3,
        "evaluations": 30,
        "reserved_usd": 0.007646058000000001,
        "state": "complete",
        "started_utc": "2026-09-22T11:01:39.111248+00:00",
        "pre_call_cumulative_evaluations": 420,
        "pre_call_reserved_usd": 0.106999368,
        "pre_call_exposure_usd": 0.106999368,
        "http_status": 200,
        "response_headers_allowlist": {
          "date": "Tue, 22 Sep 2026 11:01:38 GMT",
          "content-length": "3878",
          "content-type": "application/json"
        },
        "elapsed_seconds": 0.8398267129668966,
        "finished_utc": "2026-09-22T11:01:39.951596+00:00",
        "response_sha256": "5444a96c50a72148f38e9cdcec22fd03c70ca3ac366607ba271c2a7c7f812911",
        "reported_usage": {
          "input_tokens": 7201,
          "output_tokens": 1773
        },
        "reported_usage_estimate_usd": 0.000302442
      },
      {
        "id": "character_design_v2_reverse-r3-b01",
        "condition": "character_design_v2_reverse",
        "repeat": 3,
        "evaluations": 30,
        "reserved_usd": 0.007638708,
        "state": "complete",
        "started_utc": "2026-09-22T11:01:39.974132+00:00",
        "pre_call_cumulative_evaluations": 450,
        "pre_call_reserved_usd": 0.114638076,
        "pre_call_exposure_usd": 0.114638076,
        "http_status": 200,
        "response_headers_allowlist": {
          "date": "Tue, 22 Sep 2026 11:01:39 GMT",
          "content-length": "3879",
          "content-type": "application/json"
        },
        "elapsed_seconds": 0.8321334390202537,
        "finished_utc": "2026-09-22T11:01:40.806919+00:00",
        "response_sha256": "dc8968b56d970b74fb0fdcdd5826a2e667faca35df4b9e4b95dca8ba5958baf3",
        "reported_usage": {
          "input_tokens": 7165,
          "output_tokens": 1773
        },
        "reported_usage_estimate_usd": 0.00030093
      }
    ],
    "invoice_cost_usd": null,
    "pricing_source": "https://typesafe.ai/blog/introducing-system-one-models-and-jev",
    "cost_note": "Jev-only usage estimate, not invoice; usage counted once per request. Conservative reserves separately reported.",
    "http_request_attempts": 15,
    "evaluation_attempts": 450,
    "reserved_usd": 0.114638076,
    "reported_input_tokens": 107790,
    "reported_output_tokens": 26595,
    "reported_usage_estimate_usd": 0.0045271800000000004,
    "completed_utc": "2026-09-22T11:01:40.809566+00:00"
  },
  "caveats": [
    "The task explicitly designs a fictional character. Five forced options can produce an answer without establishing applicability, lived experience or stable AI traits.",
    "The same20 exact earlier question/selected-answer records are creative context; no gender or appearance is assigned and no desired trait scores are supplied.",
    "Original questions/options and this wrapper/report are authored; Jev selects keys, not free-form quotations.",
    "Nonunit probability totals, tied tops and returned nonwinners remain visible. Primary values use returned choices, not argmax or probability weighting; strict sensitivity separately excludes these.",
    "Human validation of Johnson120 does not validate AI psychological measurement or fictional character completion.",
    "Three same-session repeats and a30-item reversed subset describe local repeat/order variation, not psychometric reliability or causal isolation of v1 differences.",
    "Original facet labels retain source traceability; Depression/Anxiety/Neuroticism are not AI diagnoses, Intellect is not IQ, Morality is not moral worth, Liberalism is limited value/political wording.",
    "NA was deliberately removed, not converted to neutral or zero. Old v1 scores, missingness and exclusions remain unchanged."
  ]
}
