@misc{indiciaeae5cee353f47, title = {PAGAR: Taming Reward Misalignment in Inverse Reinforcement Learning-Based Imitation Learning with Protagonist Antagonist Guided Adversarial Reward}, author = {Weichao Zhou and Wenchao Li}, year = {2024}, url = {https://arxiv.org/abs/2306.01731}, note = {Source identifier: 2306.01731} }