ray/rllib/tuned_examples/dqn/atari-dist-dqn.yaml

29 lines
827 B
YAML

atari-dist-dqn:
env:
grid_search:
- BreakoutNoFrameskip-v4
- BeamRiderNoFrameskip-v4
- QbertNoFrameskip-v4
- SpaceInvadersNoFrameskip-v4
run: DQN
config:
double_q: false
dueling: false
num_atoms: 51
noisy: false
replay_buffer_config:
type: MultiAgentReplayBuffer
capacity: 1000000
num_steps_sampled_before_learning_starts: 20000
n_step: 1
target_network_update_freq: 8000
lr: .0000625
adam_epsilon: .00015
hiddens: [512]
rollout_fragment_length: 4
train_batch_size: 32
exploration_config:
epsilon_timesteps: 200000
final_epsilon: 0.01
num_gpus: 0.2
min_sample_timesteps_per_iteration: 10000