Reinforcement Learning
ml-agents
TensorBoard
Pyramids
deep-reinforcement-learning
ML-Agents-Pyramids
Instructions to use RBadal/Pyramids with libraries, inference providers, notebooks, and local apps. Follow these links to get started.
- Libraries
- ml-agents
How to use RBadal/Pyramids with ml-agents:
mlagents-load-from-hf --repo-id="RBadal/Pyramids" --local-dir="./download: string[]s"
- Notebooks
- Google Colab
- Kaggle
| { | |
| "name": "root", | |
| "metadata": { | |
| "timer_format_version": "0.1.0", | |
| "start_time_seconds": "1784182391", | |
| "python_version": "3.10.12 (main, Jul 5 2023, 18:54:27) [GCC 11.2.0]", | |
| "command_line_arguments": "/usr/local/bin/mlagents-learn ./config/ppo/PyramidsRND.yaml --env=./training-envs-executables/linux/Pyramids/Pyramids --run-id=Pyramids Training --no-graphics", | |
| "mlagents_version": "1.2.0.dev0", | |
| "mlagents_envs_version": "1.2.0.dev0", | |
| "communication_protocol_version": "1.5.0", | |
| "pytorch_version": "2.8.0+cu128", | |
| "numpy_version": "1.23.5", | |
| "end_time_seconds": "1784182458" | |
| }, | |
| "total": 67.45951297600004, | |
| "count": 1, | |
| "self": 0.5334767630001807, | |
| "children": { | |
| "run_training.setup": { | |
| "total": 0.026787270999875545, | |
| "count": 1, | |
| "self": 0.026787270999875545 | |
| }, | |
| "TrainerController.start_learning": { | |
| "total": 66.89924894199999, | |
| "count": 1, | |
| "self": 0.045259533004355035, | |
| "children": { | |
| "TrainerController._reset_env": { | |
| "total": 2.3855866989999868, | |
| "count": 1, | |
| "self": 2.3855866989999868 | |
| }, | |
| "TrainerController.advance": { | |
| "total": 64.46699979399546, | |
| "count": 1895, | |
| "self": 0.0468840670012014, | |
| "children": { | |
| "env_step": { | |
| "total": 43.18142378899688, | |
| "count": 1895, | |
| "self": 37.790354616998684, | |
| "children": { | |
| "SubprocessEnvManager._take_step": { | |
| "total": 5.364962005000962, | |
| "count": 1895, | |
| "self": 0.15036577000046236, | |
| "children": { | |
| "TorchPolicy.evaluate": { | |
| "total": 5.214596235000499, | |
| "count": 1894, | |
| "self": 5.214596235000499 | |
| } | |
| } | |
| }, | |
| "workers": { | |
| "total": 0.026107166997235254, | |
| "count": 1894, | |
| "self": 0.0, | |
| "children": { | |
| "worker_root": { | |
| "total": 66.74585728099851, | |
| "count": 1894, | |
| "is_parallel": true, | |
| "self": 32.77797824600316, | |
| "children": { | |
| "run_training.setup": { | |
| "total": 0.0, | |
| "count": 0, | |
| "is_parallel": true, | |
| "self": 0.0, | |
| "children": { | |
| "steps_from_proto": { | |
| "total": 0.002037278999978298, | |
| "count": 1, | |
| "is_parallel": true, | |
| "self": 0.0006557429999247688, | |
| "children": { | |
| "_process_rank_one_or_two_observation": { | |
| "total": 0.0013815360000535293, | |
| "count": 8, | |
| "is_parallel": true, | |
| "self": 0.0013815360000535293 | |
| } | |
| } | |
| }, | |
| "UnityEnvironment.step": { | |
| "total": 0.058151194999936706, | |
| "count": 1, | |
| "is_parallel": true, | |
| "self": 0.0005699870000626106, | |
| "children": { | |
| "UnityEnvironment._generate_step_input": { | |
| "total": 0.0004660340000555152, | |
| "count": 1, | |
| "is_parallel": true, | |
| "self": 0.0004660340000555152 | |
| }, | |
| "communicator.exchange": { | |
| "total": 0.05532181699982175, | |
| "count": 1, | |
| "is_parallel": true, | |
| "self": 0.05532181699982175 | |
| }, | |
| "steps_from_proto": { | |
| "total": 0.0017933569999968313, | |
| "count": 1, | |
| "is_parallel": true, | |
| "self": 0.00038727500009372307, | |
| "children": { | |
| "_process_rank_one_or_two_observation": { | |
| "total": 0.0014060819999031082, | |
| "count": 8, | |
| "is_parallel": true, | |
| "self": 0.0014060819999031082 | |
| } | |
| } | |
| } | |
| } | |
| } | |
| } | |
| }, | |
| "UnityEnvironment.step": { | |
| "total": 33.96787903499535, | |
| "count": 1893, | |
| "is_parallel": true, | |
| "self": 1.1110322249926412, | |
| "children": { | |
| "UnityEnvironment._generate_step_input": { | |
| "total": 0.7720298910023757, | |
| "count": 1893, | |
| "is_parallel": true, | |
| "self": 0.7720298910023757 | |
| }, | |
| "communicator.exchange": { | |
| "total": 28.453639058997396, | |
| "count": 1893, | |
| "is_parallel": true, | |
| "self": 28.453639058997396 | |
| }, | |
| "steps_from_proto": { | |
| "total": 3.6311778600029356, | |
| "count": 1893, | |
| "is_parallel": true, | |
| "self": 0.7587519300086569, | |
| "children": { | |
| "_process_rank_one_or_two_observation": { | |
| "total": 2.8724259299942787, | |
| "count": 15144, | |
| "is_parallel": true, | |
| "self": 2.8724259299942787 | |
| } | |
| } | |
| } | |
| } | |
| } | |
| } | |
| } | |
| } | |
| } | |
| } | |
| }, | |
| "trainer_advance": { | |
| "total": 21.238691937997373, | |
| "count": 1894, | |
| "self": 0.05310751299361982, | |
| "children": { | |
| "process_trajectory": { | |
| "total": 3.172604966003746, | |
| "count": 1894, | |
| "self": 3.172604966003746 | |
| }, | |
| "_update_policy": { | |
| "total": 18.012979459000007, | |
| "count": 8, | |
| "self": 9.923771062004107, | |
| "children": { | |
| "TorchPPOOptimizer.update": { | |
| "total": 8.0892083969959, | |
| "count": 663, | |
| "self": 8.0892083969959 | |
| } | |
| } | |
| } | |
| } | |
| } | |
| } | |
| }, | |
| "trainer_threads": { | |
| "total": 1.742999984344351e-06, | |
| "count": 1, | |
| "self": 1.742999984344351e-06 | |
| }, | |
| "TrainerController._save_models": { | |
| "total": 0.0014011730002039258, | |
| "count": 1, | |
| "self": 3.9758000184519915e-05, | |
| "children": { | |
| "RLTrainer._checkpoint": { | |
| "total": 0.0013614150000194059, | |
| "count": 1, | |
| "self": 0.0013614150000194059 | |
| } | |
| } | |
| } | |
| } | |
| } | |
| } | |
| } |