@misc{indiciae988ba606df02, title = {Scene-VLM: Multimodal Video Scene Segmentation via Vision-Language Models}, author = {Nimrod Berman and Adam Botach and Emanuel Ben-Baruch and Shunit Haviv Hakimi and Asaf Gendler and Ilan Naiman and Erez Yosef and Igor Kviatkovsky}, year = {2026}, url = {https://arxiv.org/abs/2512.21778}, note = {Source identifier: 2512.21778} }