@misc{indiciae7ec99836a854, title = {YUBI-STAG: Contact and Semantic-Rich Alignment for VLAs via Automated Video-Language Grounding}, author = {Masatoshi Tateno and Takehiko Ohkawa and Yueh-Hua Wu and Hanlong Li and Tatsuya Matsushima and Yoichi Sato and Kei Ota}, year = {2026}, url = {https://arxiv.org/abs/2610.09718}, note = {Source identifier: 2610.09718} }