@misc{indiciae65ab60134671, title = {On Learning Intrinsic Rewards for Policy Gradient Methods}, author = {Zeyu Zheng and Junhyuk Oh and Satinder Singh}, year = {2018}, url = {https://arxiv.org/abs/1804.06459}, note = {Source identifier: 1804.06459} }