{"true_reward": 3.455532985144788, "expected_reward": 3.4409714123681923, "runtime": 0.01867260502199999, "actions": 5.0, "seed": 7499.5, "steps": 10, "exploration_coeff": 0.5, "rollout_depth": 0}