A2C playing HalfCheetahBulletEnv-v0 from https://github.com/sgoodfriend/rl-algo-impls/tree/0760ef7d52b17f30219a27c18ba52c8895025ae3
3d6ce6f
source benchmarks/train_loop.sh | |
ALGOS="ppo" | |
ENVS="CarRacing-v0" | |
BENCHMARK_MAX_PROCS="${BENCHMARK_MAX_PROCS:-3}" | |
train_loop $ALGOS "$ENVS" | xargs -I CMD -P $BENCHMARK_MAX_PROCS bash -c CMD |