{
  "agents4sci_v2": {
    "rigor_traceability": 24.0,
    "integration_causality": 23.0,
    "feasibility_minimality": 18.0,
    "uncertainty_adaptation": 14.0,
    "decisionability": 13.0,
    "overall": 92.0
  },
  "baseline_single": {
    "rigor_traceability": 18.0,
    "integration_causality": 20.0,
    "feasibility_minimality": 15.0,
    "uncertainty_adaptation": 10.0,
    "decisionability": 5.0,
    "overall": 68.0
  },
  "baseline_tree": {
    "rigor_traceability": 22.0,
    "integration_causality": 21.0,
    "feasibility_minimality": 17.0,
    "uncertainty_adaptation": 12.0,
    "decisionability": 11.0,
    "overall": 83.0
  },
  "baseline_debate": {
    "rigor_traceability": 21.0,
    "integration_causality": 22.0,
    "feasibility_minimality": 15.0,
    "uncertainty_adaptation": 12.0,
    "decisionability": 10.0,
    "overall": 80.0
  }
}