Push agent to the Hub
Browse files- .gitattributes +1 -0
- README.md +6 -6
- logs/events.out.tfevents.1739289745.ac2833920c19.7052.4 +3 -0
- model.pt +2 -2
- replay.mp4 +0 -0
- results.json +1 -1
.gitattributes
CHANGED
|
@@ -33,3 +33,4 @@ saved_model/**/* filter=lfs diff=lfs merge=lfs -text
|
|
| 33 |
*.zip filter=lfs diff=lfs merge=lfs -text
|
| 34 |
*.zst filter=lfs diff=lfs merge=lfs -text
|
| 35 |
*tfevents* filter=lfs diff=lfs merge=lfs -text
|
|
|
|
|
|
| 33 |
*.zip filter=lfs diff=lfs merge=lfs -text
|
| 34 |
*.zst filter=lfs diff=lfs merge=lfs -text
|
| 35 |
*tfevents* filter=lfs diff=lfs merge=lfs -text
|
| 36 |
+
replay.mp4 filter=lfs diff=lfs merge=lfs -text
|
README.md
CHANGED
|
@@ -1,6 +1,6 @@
|
|
| 1 |
---
|
| 2 |
tags:
|
| 3 |
-
-
|
| 4 |
- ppo
|
| 5 |
- deep-reinforcement-learning
|
| 6 |
- reinforcement-learning
|
|
@@ -13,18 +13,18 @@ model-index:
|
|
| 13 |
type: reinforcement-learning
|
| 14 |
name: reinforcement-learning
|
| 15 |
dataset:
|
| 16 |
-
name:
|
| 17 |
-
type:
|
| 18 |
metrics:
|
| 19 |
- type: mean_reward
|
| 20 |
-
value:
|
| 21 |
name: mean_reward
|
| 22 |
verified: false
|
| 23 |
---
|
| 24 |
|
| 25 |
-
# PPO Agent Playing
|
| 26 |
|
| 27 |
-
This is a trained model of a PPO agent playing
|
| 28 |
|
| 29 |
# Hyperparameters
|
| 30 |
|
|
|
|
| 1 |
---
|
| 2 |
tags:
|
| 3 |
+
- LunarLander-v2
|
| 4 |
- ppo
|
| 5 |
- deep-reinforcement-learning
|
| 6 |
- reinforcement-learning
|
|
|
|
| 13 |
type: reinforcement-learning
|
| 14 |
name: reinforcement-learning
|
| 15 |
dataset:
|
| 16 |
+
name: LunarLander-v2
|
| 17 |
+
type: LunarLander-v2
|
| 18 |
metrics:
|
| 19 |
- type: mean_reward
|
| 20 |
+
value: 128.77 +/- 53.16
|
| 21 |
name: mean_reward
|
| 22 |
verified: false
|
| 23 |
---
|
| 24 |
|
| 25 |
+
# PPO Agent Playing LunarLander-v2
|
| 26 |
|
| 27 |
+
This is a trained model of a PPO agent playing LunarLander-v2.
|
| 28 |
|
| 29 |
# Hyperparameters
|
| 30 |
|
logs/events.out.tfevents.1739289745.ac2833920c19.7052.4
ADDED
|
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
|
|
|
|
| 1 |
+
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:c7fe792ef38d6459ad35bb8315726fe3a7b24cd63297c6b9dd266967bc07b145
|
| 3 |
+
size 306485
|
model.pt
CHANGED
|
@@ -1,3 +1,3 @@
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
-
oid sha256:
|
| 3 |
-
size
|
|
|
|
| 1 |
version https://git-lfs.github.com/spec/v1
|
| 2 |
+
oid sha256:3549e716c64df0ebb5c78afe79f4cd5c318f2a0e5ee354264180a96880bd8b9e
|
| 3 |
+
size 43026
|
replay.mp4
CHANGED
|
Binary files a/replay.mp4 and b/replay.mp4 differ
|
|
|
results.json
CHANGED
|
@@ -1 +1 @@
|
|
| 1 |
-
{"env_id": "
|
|
|
|
| 1 |
+
{"env_id": "LunarLander-v2", "mean_reward": 128.77097205634465, "std_reward": 53.15600613594416, "n_evaluation_episodes": 10, "eval_datetime": "2025-02-11T16:24:52.365171"}
|