- Wang, J., Cherian, A., "Learning Discriminative Video Representations Using Adversarial Perturbations", European Conference on Computer Vision (ECCV), September 2018.
BibTeX TR2018-139 PDF Software- @inproceedings{Wang2018sep3,
- author = {Wang, Jue and Cherian, Anoop},
- title = {{Learning Discriminative Video Representations Using Adversarial Perturbations}},
- booktitle = {European Conference on Computer Vision (ECCV)},
- year = 2018,
- month = sep,
- url = {https://www.merl.com/publications/TR2018-139}
- }
- Kocanaogullari, A., Ataer-Cansizoglu, E., "Active Descriptor Learning for Feature Matching", International Workshop on Compact and Efficient Feature Representation and Learning in Computer Vision, September 2018.
BibTeX TR2018-132 PDF- @inproceedings{Kocanaogullari2018sep,
- author = {Kocanaogullari, Aziz and Ataer-Cansizoglu, Esra},
- title = {{Active Descriptor Learning for Feature Matching}},
- booktitle = {International Workshop on Compact and Efficient Feature Representation and Learning in Computer Vision},
- year = 2018,
- month = sep,
- url = {https://www.merl.com/publications/TR2018-132}
- }
- Jones, M.J., Broad, A., Lee, T.-Y., "Recurrent Multi-frame Single Shot Detector for Video Object Detection", British Machine Vision Conference (BMVC), September 2018.
BibTeX TR2018-137 PDF- @inproceedings{Jones2018sep,
- author = {Jones, Michael J. and Broad, Alexander and Lee, Teng-Yok},
- title = {{Recurrent Multi-frame Single Shot Detector for Video Object Detection}},
- booktitle = {British Machine Vision Conference (BMVC)},
- year = 2018,
- month = sep,
- url = {https://www.merl.com/publications/TR2018-137}
- }
- Wang, Z.-Q., Le Roux, J., Wang, D., Hershey, J., "End-to-End Speech Separation with Unfolded Iterative Phase Reconstruction", Interspeech, September 2018.
BibTeX TR2018-135 PDF- @inproceedings{Wang2018sep,
- author = {Wang, Zhong-Qiu and {Le Roux}, Jonathan and Wang, DeLiang and Hershey, John},
- title = {{End-to-End Speech Separation with Unfolded Iterative Phase Reconstruction}},
- booktitle = {Interspeech},
- year = 2018,
- month = sep,
- url = {https://www.merl.com/publications/TR2018-135}
- }
- Watanabe, S., Hori, T., Karita, S., Hayashi, T., Nishitoba, J., Unno, Y., Enrique Yalta Soplin, N., Heymann, J., Wiesner, M., Chen, N., Renduchintala, A., Ochiai, T., "ESPnet: End-to-End Speech Processing Toolkit", Interspeech, September 2018.
BibTeX TR2018-136 PDF- @inproceedings{Watanabe2018sep,
- author = {Watanabe, Shinji and Hori, Takaaki and Karita, Shigeki and Hayashi, Tomoki and Nishitoba, Jiro and Unno, Yuya and Enrique Yalta Soplin, Nelson and Heymann, Jahn and Wiesner, Matthew and Chen, Nanxin and Renduchintala, Adithya and Ochiai, Tsubasa},
- title = {{ESPnet: End-to-End Speech Processing Toolkit}},
- booktitle = {Interspeech},
- year = 2018,
- month = sep,
- url = {https://www.merl.com/publications/TR2018-136}
- }
- Ataer-Cansizoglu, E., Jones, M.J., Zhang, Z., Sullivan, A., "Verification of Very Low-Resolution Faces Using An Identity-Preserving Deep Face Super-resolution Network", arXiv, August 2018.
BibTeX arXiv- @article{Ataer-Cansizoglu2018aug,
- author = {Ataer-Cansizoglu, Esra and Jones, Michael J. and Zhang, Ziming and Sullivan, Alan},
- title = {{Verification of Very Low-Resolution Faces Using An Identity-Preserving Deep Face Super-resolution Network}},
- journal = {arXiv},
- year = 2018,
- month = aug,
- url = {https://arxiv.org/abs/1903.10974}
- }
- Zhang, Z., "LMKL-Net: A Fast Localized Multiple Kernel Learning Solver via Deep Neural Networks", arXiv, July 12, 2018.
BibTeX arXiv- @article{Zhang2018jul,
- author = {Zhang, Ziming},
- title = {{LMKL-Net: A Fast Localized Multiple Kernel Learning Solver via Deep Neural Networks}},
- journal = {arXiv},
- year = 2018,
- month = jul,
- url = {https://arxiv.org/abs/1805.08656}
- }
- Pan, Y., Farahmand, A.-M., White, M., Nabi, S., Grover, P., Nikovski, D.N., "Reinforcement Learning with Function-Valued Action Spaces for Partial Differential Equation Control", International Conference on Machine Learning (ICML), July 2018.
BibTeX TR2018-101 PDF- @inproceedings{Pan2018jul,
- author = {Pan, Yangchen and Farahmand, Amir-massoud and White, Martha and Nabi, Saleh and Grover, Piyush and Nikovski, Daniel N.},
- title = {{Reinforcement Learning with Function-Valued Action Spaces for Partial Differential Equation Control}},
- booktitle = {International Conference on Machine Learning (ICML)},
- year = 2018,
- month = jul,
- url = {https://www.merl.com/publications/TR2018-101}
- }
- Barker, J., Marxer, R., Vincent, E., Watanabe, S., "The CHiME challenges: Robust speech recognition in everyday environments" in New Era for Robust Speech Recognition: Exploiting Deep Learning, Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R., Eds., chapter 14, Springer, July 2018.
BibTeX - @incollection{Barker2018jul,
- author = {Barker, Jon and Marxer, Ricard and Vincent, Emmanuel and Watanabe, Shinji},
- title = {{The CHiME challenges: Robust speech recognition in everyday environments}},
- booktitle = {New Era for Robust Speech Recognition: Exploiting Deep Learning},
- year = 2018,
- editor = {Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R.},
- chapter = 14,
- month = jul,
- publisher = {Springer}
- }
- Erdogan, H., Hershey, J., Watanabe, S., Le Roux, J., "Deep recurrent networks for separation and recognition of single-channel speech in non-stationary background audio" in New Era for Robust Speech Recognition: Exploiting Deep Learning, Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R., Eds., chapter 7, Springer, July 2018.
BibTeX - @incollection{Erdogan2018jul,
- author = {Erdogan, Hakan and Hershey, John and Watanabe, Shinji and {Le Roux}, Jonathan},
- title = {{Deep recurrent networks for separation and recognition of single-channel speech in non-stationary background audio}},
- booktitle = {New Era for Robust Speech Recognition: Exploiting Deep Learning},
- year = 2018,
- editor = {Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R.},
- chapter = 7,
- month = jul,
- publisher = {Springer},
- isbn = {978-3-319-64680-0}
- }
- Hershey, J., Le Roux, J., Watanabe, S., Wisdom, S., Chen, Z., Isik, Y., "Novel deep architectures in speech processing" in New Era for Robust Speech Recognition: Exploiting Deep Learning, Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R., Eds., chapter 6, Springer, July 9, 2018.
BibTeX - @incollection{Hershey2018jul,
- author = {Hershey, John and {Le Roux}, Jonathan and Watanabe, Shinji and Wisdom, Scott and Chen, Zhuo and Isik, Yusuf},
- title = {{Novel deep architectures in speech processing}},
- booktitle = {New Era for Robust Speech Recognition: Exploiting Deep Learning},
- year = 2018,
- editor = {Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R.},
- chapter = 6,
- month = jul,
- publisher = {Springer}
- }
- Karafiat, M., Vesely, K., Zmolikova, K., Delcroix, M., Watanabe, S., Burget, L., Cernocky, J., Szoke, I., Novotny, O., "Training data augmentation and data selectio" in New Era for Robust Speech Recognition: Exploiting Deep Learning, Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R., Eds., chapter 10, Springer, July 9, 2018.
BibTeX - @incollection{Karafiat2018jul,
- author = {Karafiat, Martin and Vesely, Karel and Zmolikova, Katerina and Delcroix, Marc and Watanabe, Shinji and Burget, Lukas and Cernocky, Jan and Szoke, Igor and Novotny, Ondrej},
- title = {{Training data augmentation and data selectio}},
- booktitle = {New Era for Robust Speech Recognition: Exploiting Deep Learning},
- year = 2018,
- editor = {Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R.},
- chapter = 10,
- month = jul,
- publisher = {Springer}
- }
- Watanabe, S., Hori, T., Miao, Y., Delcroix, M., Metze, F., Hershey, J., "Toolkits for robust speech processing" in New Era for Robust Speech Recognition: Exploiting Deep Learning, Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R., Eds., chapter 14, Springer, July 9, 2018.
BibTeX - @incollection{Watanabe2018jul,
- author = {Watanabe, Shinji and Hori, Takaaki and Miao, Yajie and Delcroix, Marc and Metze, Florian and Hershey, John},
- title = {{Toolkits for robust speech processing}},
- booktitle = {New Era for Robust Speech Recognition: Exploiting Deep Learning},
- year = 2018,
- editor = {Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R.},
- chapter = 14,
- month = jul,
- publisher = {Springer}
- }
- Xiao, X., Watanabe, S., Erdogan, H., Mandel, M., Lu, L., Hershey, J., Seltzer, M., Chen, G., Zhang, Y., Yu, D., "Discriminative beamforming with phase aware neural networks for speech enhancement and recognition" in New Era for Robust Speech Recognition: Exploiting Deep Learning, Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R., Eds., chapter 4, Springer, July 9, 2018.
BibTeX - @incollection{Xiao2018jul2,
- author = {Xiao, Xiong and Watanabe, Shinji and Erdogan, Hakan and Mandel, Michael and Lu, Liang and Hershey, John and Seltzer, Mike and Chen, Guoguo and Zhang, Yu and Yu, Dong},
- title = {{Discriminative beamforming with phase aware neural networks for speech enhancement and recognition}},
- booktitle = {New Era for Robust Speech Recognition: Exploiting Deep Learning},
- year = 2018,
- editor = {Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R.},
- chapter = 4,
- month = jul,
- publisher = {Springer}
- }
- Wang, P., Fu, H., Koike-Akino, T., Orlik, P.V., "Multi-Layer Terahertz Imaging of Non-Overlapping Contents", IEEE Sensor Array and Multi-Channel Signal Processing Workshop (IEEE SAM), DOI: 10.1109/SAM.2018.8448438, July 2018, pp. 652-656.
BibTeX TR2018-098 PDF- @inproceedings{Wang2018jul2,
- author = {Wang, Pu and Fu, Haoyu and Koike-Akino, Toshiaki and Orlik, Philip V.},
- title = {{Multi-Layer Terahertz Imaging of Non-Overlapping Contents}},
- booktitle = {IEEE Sensor Array and Multi-Channel Signal Processing Workshop (IEEE SAM)},
- year = 2018,
- pages = {652--656},
- month = jul,
- doi = {10.1109/SAM.2018.8448438},
- url = {https://www.merl.com/publications/TR2018-098}
- }
- Liu, J., Guo, J., Orlik, P.V., Shibata, M., Nakahara, D., Mii, S., Takac, M., "Anomaly Detection in Manufacturing Systems Using Structured Neural Networks", IEEE World Congress on Intelligent Control and Automation, DOI: 10.1109/WCICA.2018.8630692, July 2018, pp. 175-180.
BibTeX TR2018-097 PDF- @inproceedings{Liu2018jul2,
- author = {Liu, Jie and Guo, Jianlin and Orlik, Philip V. and Shibata, Masahiko and Nakahara, Daiki and Mii, Satoshi and Takac, Martin},
- title = {{Anomaly Detection in Manufacturing Systems Using Structured Neural Networks}},
- booktitle = {IEEE World Congress on Intelligent Control and Automation},
- year = 2018,
- pages = {175--180},
- month = jul,
- doi = {10.1109/WCICA.2018.8630692},
- url = {https://www.merl.com/publications/TR2018-097}
- }
- Koike-Akino, T., Millar, D.S., Parsons, K., Kojima, K., "Fiber Nonlinearity Equalization with Multi-Label Deep Learning Scalable to High-Order DP-QAM", Signal Processing in Photonic Communications (SPPCom), DOI: 10.1364/SPPCOM.2018.SpM4G.1, July 2018.
BibTeX TR2018-047 PDF- @inproceedings{Koike-Akino2018jul3,
- author = {Koike-Akino, Toshiaki and Millar, David S. and Parsons, Kieran and Kojima, Keisuke},
- title = {{Fiber Nonlinearity Equalization with Multi-Label Deep Learning Scalable to High-Order DP-QAM}},
- booktitle = {Signal Processing in Photonic Communications (SPPCom)},
- year = 2018,
- month = jul,
- doi = {10.1364/SPPCOM.2018.SpM4G.1},
- url = {https://www.merl.com/publications/TR2018-047}
- }
- Shen, Y., Feng, C., Yang, Y., Tian, D., "Mining Point Cloud Local Structures by Kernel Correlation and Graph Pooling", IEEE Conference on Computer Vision and Pattern Recognition (CVPR), June 2018.
BibTeX TR2018-041 PDF Software- @inproceedings{Shen2018jun,
- author = {Shen, Yiru and Feng, Chen and Yang, Yaoqing and Tian, Dong},
- title = {{Mining Point Cloud Local Structures by Kernel Correlation and Graph Pooling}},
- booktitle = {IEEE Conference on Computer Vision and Pattern Recognition (CVPR)},
- year = 2018,
- month = jun,
- url = {https://www.merl.com/publications/TR2018-041}
- }
- Yang, Y., Feng, C., Shen, Y., Tian, D., "FoldingNet: Point Cloud Auto-encoder via Deep Grid Deformation", IEEE Conference on Computer Vision and Pattern Recognition (CVPR), DOI: 10.1109/CVPR.2018.00029, June 2018.
BibTeX TR2018-042 PDF Video Software- @inproceedings{Yang2018jun,
- author = {Yang, Yaoqing and Feng, Chen and Shen, Yiru and Tian, Dong},
- title = {{FoldingNet: Point Cloud Auto-encoder via Deep Grid Deformation}},
- booktitle = {IEEE Conference on Computer Vision and Pattern Recognition (CVPR)},
- year = 2018,
- month = jun,
- doi = {10.1109/CVPR.2018.00029},
- url = {https://www.merl.com/publications/TR2018-042}
- }
- Zhang, Z., Wu, Y., Wang, G., "BPGrad: Towards Global Optimality in Deep Learning via Branch and Pruning", IEEE Conference on Computer Vision and Pattern Recognition (CVPR), June 2018, pp. 3301-3309.
BibTeX TR2018-068 PDF- @inproceedings{Zhang2018jun,
- author = {Zhang, Ziming and Wu, Yuanwei and Wang, Guanghui},
- title = {{BPGrad: Towards Global Optimality in Deep Learning via Branch and Pruning}},
- booktitle = {IEEE Conference on Computer Vision and Pattern Recognition (CVPR)},
- year = 2018,
- pages = {3301--3309},
- month = jun,
- url = {https://www.merl.com/publications/TR2018-068}
- }
- Fujihashi, T., Koike-Akino, T., Watanabe, T., Orlik, P.V., "Nonlinear Equalization with Deep Learning for Multi-Purpose Visual MIMO Communications", IEEE International Conference on Communications (ICC), DOI: 10.1109/ICC.2018.8422544, May 2018.
BibTeX TR2018-039 PDF- @inproceedings{Fujihashi2018may,
- author = {Fujihashi, Takuya and Koike-Akino, Toshiaki and Watanabe, Takashi and Orlik, Philip V.},
- title = {{Nonlinear Equalization with Deep Learning for Multi-Purpose Visual MIMO Communications}},
- booktitle = {IEEE International Conference on Communications (ICC)},
- year = 2018,
- month = may,
- doi = {10.1109/ICC.2018.8422544},
- url = {https://www.merl.com/publications/TR2018-039}
- }
- Ochiai, T., Watanabe, S., Katagiri, S., Hori, T., Hershey, J.R., "Speaker Adaptation for Multichannel End-to-End Speech Recognition", IEEE International Conference on Acoustics, Speech, and Signal Processing (ICASSP), DOI: 10.1109/ICASSP.2018.8462161, April 2018, pp. 6707-6711.
BibTeX TR2018-006 PDF- @inproceedings{Ochiai2018apr,
- author = {Ochiai, Tsubasa and Watanabe, Shinji and Katagiri, Shigeru and Hori, Takaaki and Hershey, John R.},
- title = {{Speaker Adaptation for Multichannel End-to-End Speech Recognition}},
- booktitle = {IEEE International Conference on Acoustics, Speech, and Signal Processing (ICASSP)},
- year = 2018,
- pages = {6707--6711},
- month = apr,
- doi = {10.1109/ICASSP.2018.8462161},
- url = {https://www.merl.com/publications/TR2018-006}
- }
- Settle, S., Le Roux, J., Hori, T., Watanabe, S., Hershey, J.R., "End-to-End Multi-Speaker Speech Recognition", IEEE International Conference on Acoustics, Speech, and Signal Processing (ICASSP), DOI: 10.1109/ICASSP.2018.8461893, April 2018, pp. 4819-4823.
BibTeX TR2018-001 PDF Video- @inproceedings{Settle2018apr,
- author = {Settle, Shane and {Le Roux}, Jonathan and Hori, Takaaki and Watanabe, Shinji and Hershey, John R.},
- title = {{End-to-End Multi-Speaker Speech Recognition}},
- booktitle = {IEEE International Conference on Acoustics, Speech, and Signal Processing (ICASSP)},
- year = 2018,
- pages = {4819--4823},
- month = apr,
- doi = {10.1109/ICASSP.2018.8461893},
- url = {https://www.merl.com/publications/TR2018-001}
- }
- Wang, Z.-Q., Le Roux, J., Hershey, J.R., "Multi-Channel Deep Clustering: Discriminative Spectral and Spatial Embeddings for Speaker-Independent Speech Separation", IEEE International Conference on Acoustics, Speech, and Signal Processing (ICASSP), DOI: 10.1109/ICASSP.2018.8461639, April 2018, pp. 1-5.
BibTeX TR2018-007 PDF- @inproceedings{Wang2018apr2,
- author = {Wang, Zhong-Qiu and {Le Roux}, Jonathan and Hershey, John R.},
- title = {{Multi-Channel Deep Clustering: Discriminative Spectral and Spatial Embeddings for Speaker-Independent Speech Separation}},
- booktitle = {IEEE International Conference on Acoustics, Speech, and Signal Processing (ICASSP)},
- year = 2018,
- pages = {1--5},
- month = apr,
- doi = {10.1109/ICASSP.2018.8461639},
- url = {https://www.merl.com/publications/TR2018-007}
- }
- Wang, Z.-Q., Le Roux, J., Hershey, J.R., "Alternative Objective Functions for Deep Clustering", IEEE International Conference on Acoustics, Speech, and Signal Processing (ICASSP), DOI: 10.1109/ICASSP.2018.8462507, April 2018, pp. 686-690.
BibTeX TR2018-005 PDF- @inproceedings{Wang2018apr,
- author = {Wang, Zhong-Qiu and {Le Roux}, Jonathan and Hershey, John R.},
- title = {{Alternative Objective Functions for Deep Clustering}},
- booktitle = {IEEE International Conference on Acoustics, Speech, and Signal Processing (ICASSP)},
- year = 2018,
- pages = {686--690},
- month = apr,
- doi = {10.1109/ICASSP.2018.8462507},
- url = {https://www.merl.com/publications/TR2018-005}
- }