actor_bc_coef: 0.05
actor_learning_rate: 0.001
actor_ln: false
actor_n_hiddens: 3
batch_size: 1024
critic_bc_coef: 0.5
critic_learning_rate: 0.001
critic_ln: true
critic_n_hiddens: 3
dataset_name: hopper-medium-replay-v2
eval_episodes: 10
eval_every: 10
eval_seed: 42
gamma: 0.99
group: rebrac-hopper-medium-replay-v2
hidden_dim: 256
name: rebrac
noise_clip: 0.5
normalize_q: true
normalize_reward: false
normalize_states: false
num_epochs: 1000
num_updates_on_epoch: 1000
policy_freq: 2
policy_noise: 0.2
project: ReBRAC
tau: 0.005
train_seed: 0
n_classes: 101
sigma_frac: 0.75
