{"mean_reward": 589.0, "std_reward": 149.56269588369955, "is_deterministic": false, "n_eval_episodes": 10, "eval_datetime": "2023-12-07T15:19:09.041166"}