@misc{indiciaeb18183c5fe09, title = {SPLASH! Sample-efficient Preference-based inverse reinforcement learning for Long-horizon Adversarial tasks from Suboptimal Hierarchical demonstrations}, author = {Peter Crowley and Zachary Serlin and Tyler Paine and Makai Mann and Michael Benjamin and Calin Belta}, year = {2025}, url = {https://arxiv.org/abs/2507.08707}, note = {Source identifier: 2507.08707} }