@misc{indiciaeb1260ddc3d77, title = {Reward-Adaptive Reinforcement Learning: Dynamic Policy Gradient Optimization for Bipedal Locomotion}, author = {Changxin Huang and Guangrun Wang and Zhibo Zhou and Ronghui Zhang and Liang Lin}, year = {2021}, url = {https://arxiv.org/abs/2107.01908}, note = {Source identifier: 2107.01908} }