@misc{indiciae27e7a343040b, title = {Reliable Off-policy Evaluation for Reinforcement Learning}, author = {Jie Wang and Rui Gao and Hongyuan Zha}, year = {2022}, doi = {10.1287/opre.2022.2382}, url = {https://arxiv.org/abs/2011.04102}, note = {Source identifier: 2011.04102} }