@inproceedings{c48382e215bd48cab6533a56844dde39,
title = "Spatio-temporal multi-scale soft quantization learning for skeleton-based human action recognition",
abstract = "Effective feature representation is important for action recognition. In this paper, a novel soft quantization learning method is proposed to represent visual features for action recognition. Specifically, we propose a dual multi-scale soft-quantization network, which is a trainable quantizer using RBF neurons. The RBF layer includes dual multi-scale structure, namely a three-level hierarchical skeleton structure in space, and a temporal-pyramid based multi-scale time structure. Different spatial levels in the RBF layer have respective RBF neurons for hierarchical spatial information, while the temporal scales share them to reduce the number of parameters in the network. An accumulation layer following the RBF layer summarizes the RBF output as a histogram representation for classification task. The proposed method is end-to-end differentiable that can be trained using regular back-propagation. The conducted experiments on benchmark datasets verify that the proposed method outperforms state-of-the-art methods.",
keywords = "Action recognition, Bag-of-features, Multi-scale, Soft quantization",
author = "Jianyu Yang and Chen Zhu and Junsong Yuan",
note = "Publisher Copyright: {\textcopyright} 2019 IEEE.; 2019 IEEE International Conference on Multimedia and Expo, ICME 2019 ; Conference date: 08-07-2019 Through 12-07-2019",
year = "2019",
month = jul,
doi = "10.1109/ICME.2019.00189",
language = "English",
series = "Proceedings - IEEE International Conference on Multimedia and Expo",
publisher = "IEEE Computer Society",
pages = "1078--1083",
booktitle = "Proceedings - 2019 IEEE International Conference on Multimedia and Expo, ICME 2019",
address = "United States",
}