@misc{indiciae27ce581e6b65, title = {Stable deep reinforcement learning method by predicting uncertainty in rewards as a subtask}, author = {Kanata Suzuki and Tetsuya Ogata}, year = {2021}, doi = {10.1007/978-3-030-63833-7\_55}, url = {https://arxiv.org/abs/2101.06906}, note = {Source identifier: 2101.06906} }