@misc{indiciaecdc0e54cc3da, title = {Cross-Modal Visuo-Tactile Representation Learning with Action Chunking Transformers for Contact-Rich Manipulation}, author = {Yaohua Liu and Rong Fu and Amir H. Gandomi and Simon Fong and Hengjun Zhang}, year = {2026}, url = {https://arxiv.org/abs/2602.00514}, note = {Source identifier: 2602.00514} }