@misc{indiciaed7841223e79f, title = {Proprioception Enhances Vision Language Model in Generating Captions and Subtask Segmentations for Robot Task}, author = {Kanata Suzuki and Shota Shimizu and Tetsuya Ogata}, year = {2026}, url = {https://arxiv.org/abs/2512.20876}, note = {Source identifier: 2512.20876} }