cleth
/

poca-SoccerTwos

+---
+      tags:
+      - unity-ml-agents
+      - ml-agents
+      - deep-reinforcement-learning
+      - reinforcement-learning
+      - ML-Agents-SoccerTwos
+      library_name: ml-agents
+---
+  # **poca** Agent playing **SoccerTwos**
+  This is a trained model of a **poca** agent playing **SoccerTwos** using the [Unity ML-Agents Library](https://github.com/Unity-Technologies/ml-agents).
+  ## Usage (with ML-Agents)
+  The Documentation: https://github.com/huggingface/ml-agents#get-started
+  We wrote a complete tutorial to learn to train your first agent using ML-Agents and publish it to the Hub:
+  ### Resume the training
+  ```
+  mlagents-learn <your_configuration_file_path.yaml> --run-id=<run_id> --resume
+  ```
+  ### Watch your Agent play
+  You can watch your agent **playing directly in your browser:**.
+  1. Go to https://huggingface.co/spaces/unity/ML-Agents-SoccerTwos
+  2. Step 1: Write your model_id: cleth/poca-SoccerTwos
+  3. Step 2: Select your *.nn /*.onnx file
+  4. Click on Watch the agent play 👀

SoccerTwos.onnx ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:14294e89418f415fee570469e377bbb9d04e6ee6f207569fb0cf7db586c8406c
+size 1764633

configuration.yaml ADDED Viewed

	@@ -0,0 +1,82 @@

+default_settings: null
+behaviors:
+  SoccerTwos:
+    trainer_type: poca
+    hyperparameters:
+      batch_size: 4096
+      buffer_size: 40960
+      learning_rate: 0.0003
+      beta: 0.005
+      epsilon: 0.2
+      lambd: 0.95
+      num_epoch: 3
+      learning_rate_schedule: constant
+      beta_schedule: constant
+      epsilon_schedule: constant
+    checkpoint_interval: 500000
+    network_settings:
+      normalize: false
+      hidden_units: 512
+      num_layers: 2
+      vis_encode_type: simple
+      memory: null
+      goal_conditioning_type: hyper
+      deterministic: false
+    reward_signals:
+      extrinsic:
+        gamma: 0.99
+        strength: 1.0
+        network_settings:
+          normalize: false
+          hidden_units: 128
+          num_layers: 2
+          vis_encode_type: simple
+          memory: null
+          goal_conditioning_type: hyper
+          deterministic: false
+    init_path: null
+    keep_checkpoints: 5
+    even_checkpoints: false
+    max_steps: 5000000
+    time_horizon: 1000
+    summary_freq: 10000
+    threaded: true
+    self_play:
+      save_steps: 50000
+      team_change: 200000
+      swap_steps: 2000
+      window: 30
+      play_against_latest_model_ratio: 0.5
+      initial_elo: 1200.0
+    behavioral_cloning: null
+env_settings:
+  env_path: ./training-envs-executables/SoccerTwos/SoccerTwos.app
+  env_args: null
+  base_port: 5005
+  num_envs: 1
+  num_areas: 1
+  seed: -1
+  max_lifetime_restarts: 10
+  restarts_rate_limit_n: 1
+  restarts_rate_limit_period_s: 60
+engine_settings:
+  width: 84
+  height: 84
+  quality_level: 5
+  time_scale: 20
+  target_frame_rate: -1
+  capture_frame_rate: 60
+  no_graphics: true
+environment_parameters: null
+checkpoint_settings:
+  run_id: SoccerTwos
+  initialize_from: null
+  load_model: false
+  resume: false
+  force: true
+  train_model: false
+  inference: false
+  results_dir: results
+torch_settings:
+  device: null
+debug: false