@misc{indiciae5a2fbc3a8ba7, title = {Multiple-Step Greedy Policies in Online and Approximate Reinforcement Learning}, author = {Yonathan Efroni and Gal Dalal and Bruno Scherrer and Shie Mannor}, year = {2018}, url = {https://arxiv.org/abs/1805.07956}, note = {Source identifier: 1805.07956} }