@misc{indiciaed0045d4e66b9, title = {OPIRL: Sample Efficient Off-Policy Inverse Reinforcement Learning via Distribution Matching}, author = {Hana Hoshino and Kei Ota and Asako Kanezaki and Rio Yokota}, year = {2022}, url = {https://arxiv.org/abs/2109.04307}, note = {Source identifier: 2109.04307} }