@misc{indiciaedacb161b78da, title = {VLT: A Vision-Language-Time Series Multimodal Foundation Model for Industrial Intelligence}, author = {Haiteng Wang and Jingheng Yan and Xiaokang Wang and Lei Ren}, year = {2026}, url = {https://arxiv.org/abs/2607.14510}, note = {Source identifier: 2607.14510} }