{"mean_reward": 5.0, "std_reward": 7.0710678118654755, "is_deterministic": false, "n_eval_episodes": 10, "eval_datetime": "2023-09-26T21:18:11.661490"} |
{"mean_reward": 5.0, "std_reward": 7.0710678118654755, "is_deterministic": false, "n_eval_episodes": 10, "eval_datetime": "2023-09-26T21:18:11.661490"} |