@misc{indiciae462152f77e36, title = {Explore, Exploit or Listen: Combining Human Feedback and Policy Model to Speed up Deep Reinforcement Learning in 3D Worlds}, author = {Zhiyu Lin and Brent Harrison and Aaron Keech and Mark O. Riedl}, year = {2021}, url = {https://arxiv.org/abs/1709.03969}, note = {Source identifier: 1709.03969} }