@misc{indiciaef6147a3899cb, title = {Soft Policy Gradient Method for Maximum Entropy Deep Reinforcement Learning}, author = {Wenjie Shi and Shiji Song and Cheng Wu}, year = {2019}, url = {https://arxiv.org/abs/1909.03198}, note = {Source identifier: 1909.03198} }