@misc{indiciaea9466f9c6f01, title = {Videos are Sample-Efficient Supervisions: Behavior Cloning from Videos via Latent Representations}, author = {Xin Liu and Haoran Li and Dongbin Zhao}, year = {2025}, url = {https://arxiv.org/abs/2512.21586}, note = {Source identifier: 2512.21586} }