@misc{indiciae60342022a93b, title = {Reinforcement Online Learning to Rank with Unbiased Reward Shaping}, author = {Shengyao Zhuang and Zhihao Qiao and Guido Zuccon}, year = {2022}, url = {https://arxiv.org/abs/2201.01534}, note = {Source identifier: 2201.01534} }