@misc{indiciaeec24cab78e25, title = {Deep Reinforcement Learning from Policy-Dependent Human Feedback}, author = {Dilip Arumugam and Jun Ki Lee and Sophie Saskin and Michael L. Littman}, year = {2019}, url = {https://arxiv.org/abs/1902.04257}, note = {Source identifier: 1902.04257} }