@misc{indiciae5dc1ded07b91, title = {Segment-to-Act: Label-Noise-Robust Action-Prompted Video Segmentation Towards Embodied Intelligence}, author = {Wenxin Li and Kunyu Peng and Di Wen and Ruiping Liu and Mengfei Duan and Kai Luo and Kailun Yang}, year = {2026}, url = {https://arxiv.org/abs/2509.16677}, note = {Source identifier: 2509.16677} }