Publications

TaherSima, M., Kojima, K., Koike-Akino, T., Jha, D.K., Wang, B., Lin, C., Parsons, K., "Deep Neural Network Inverse Modeling for Integrated Photonics", Optical Fiber Communication Conference and Exposition and the National Fiber Optic Engineers Conference (OFC/NFOEC), DOI: 10.1364/OFC.2019.W3B.5, March 2019.
BibTeX TR2018-183 PDF
- @inproceedings{TaherSima2019mar,
- author = {TaherSima, Mohammad and Kojima, Keisuke and Koike-Akino, Toshiaki and Jha, Devesh K. and Wang, Bingnan and Lin, Chungwei and Parsons, Kieran},
- title = {{Deep Neural Network Inverse Modeling for Integrated Photonics}},
- booktitle = {Optical Fiber Communication Conference and Exposition and the National Fiber Optic Engineers Conference (OFC/NFOEC)},
- year = 2019,
- month = mar,
- doi = {10.1364/OFC.2019.W3B.5},
- url = {https://www.merl.com/publications/TR2018-183}
- }
TaherSima, M., Kojima, K., Koike-Akino, T., Jha, D., Wang, B., Lin, C., Parsons, K., "Deep Neural Network Inverse Design of Integrated Photonic Power Splitters", Nature Scientific Reports, DOI: 10.1038/s41598-018-37952-2, Vol. 9, pp. 1368, December 2018.
BibTeX TR2018-180 PDF
- @article{TaherSima2018dec,
- author = {TaherSima, Mohammad and Kojima, Keisuke and Koike-Akino, Toshiaki and Jha, Devesh and Wang, Bingnan and Lin, Chungwei and Parsons, Kieran},
- title = {{Deep Neural Network Inverse Design of Integrated Photonic Power Splitters}},
- journal = {Nature Scientific Reports},
- year = 2018,
- volume = 9,
- pages = 1368,
- month = dec,
- doi = {10.1038/s41598-018-37952-2},
- issn = {2045-2322},
- url = {https://www.merl.com/publications/TR2018-180}
- }
Paul, S., van Baar, J., "Trajectory-based Learning for Ball-in-Maze Games", NIPS Workshop on Imitation Learning and its Challenges in Robotics, December 2018.
BibTeX TR2018-158 PDF
- @inproceedings{Paul2018dec,
- author = {Paul, Sujoy and {van Baar}, Jeroen},
- title = {{Trajectory-based Learning for Ball-in-Maze Games}},
- booktitle = {NIPS Workshop on Imitation Learning and its Challenges in Robotics},
- year = 2018,
- month = dec,
- url = {https://www.merl.com/publications/TR2018-158}
- }
Jha, D.K., Romeres, D., van Baar, J., Sullivan, A., Nikovski, D.N., "Learning Tasks in a Complex Circular Maze Environment", NIPS Workshop on Modeling the Physical World: Perception, Learning, and Control, December 2018.
BibTeX TR2018-169 PDF
- @inproceedings{vanBaar2018dec,
- author = {Jha, Devesh K. and Romeres, Diego and {van Baar}, Jeroen and Sullivan, Alan and Nikovski, Daniel N.},
- title = {{Learning Tasks in a Complex Circular Maze Environment}},
- booktitle = {NIPS Workshop on Modeling the Physical World: Perception, Learning, and Control},
- year = 2018,
- month = dec,
- url = {https://www.merl.com/publications/TR2018-169}
- }
Yu, X., Chaturvedi, S., Feng, C., Taguchi, Y., Lee, T.-Y., Fernandes, C., Ramalingam, S., "VLASE: Vehicle Localization by Aggregating Semantic Edges", IEEE/RSJ International Conference on Intelligent Robots and Systems (IROS), DOI: 10.1109/IROS.2018.8594358, October 2018, pp. 3196-3203.
BibTeX TR2018-113 PDF
- @inproceedings{Yu2018oct,
- author = {Yu, Xin and Chaturvedi, Sagar and Feng, Chen and Taguchi, Yuichi and Lee, Teng-Yok and Fernandes, Clinton and Ramalingam, Srikumar},
- title = {{VLASE: Vehicle Localization by Aggregating Semantic Edges}},
- booktitle = {IEEE/RSJ International Conference on Intelligent Robots and Systems (IROS)},
- year = 2018,
- pages = {3196--3203},
- month = oct,
- doi = {10.1109/IROS.2018.8594358},
- url = {https://www.merl.com/publications/TR2018-113}
- }
Ataer-Cansizoglu, E., Jones, M.J., "Super-resolution of Very Low-Resolution Faces from Videos", British Machine Vision Conference (BMVC), September 2018.
BibTeX TR2018-140 PDF
- @inproceedings{Ataer-Cansizoglu2018sep,
- author = {Ataer-Cansizoglu, Esra and Jones, Michael J.},
- title = {{Super-resolution of Very Low-Resolution Faces from Videos}},
- booktitle = {British Machine Vision Conference (BMVC)},
- year = 2018,
- month = sep,
- url = {https://www.merl.com/publications/TR2018-140}
- }
Wichern, G., Le Roux, J., "Phase Reconstruction with Learned Time-Frequency Representations for Single-Channel Speech Separation", International Workshop on Acoustic Signal Enhancement (IWAENC), DOI: 10.1109/IWAENC.2018.8521243, September 2018.
BibTeX TR2018-146 PDF
- @inproceedings{Wichern2018sep,
- author = {Wichern, Gordon and {Le Roux}, Jonathan},
- title = {{Phase Reconstruction with Learned Time-Frequency Representations for Single-Channel Speech Separation}},
- booktitle = {International Workshop on Acoustic Signal Enhancement (IWAENC)},
- year = 2018,
- month = sep,
- doi = {10.1109/IWAENC.2018.8521243},
- url = {https://www.merl.com/publications/TR2018-146}
- }
Wang, J., Cherian, A., "Learning Discriminative Video Representations Using Adversarial Perturbations", European Conference on Computer Vision (ECCV), September 2018.
BibTeX TR2018-139 PDF Software
- @inproceedings{Wang2018sep3,
- author = {Wang, Jue and Cherian, Anoop},
- title = {{Learning Discriminative Video Representations Using Adversarial Perturbations}},
- booktitle = {European Conference on Computer Vision (ECCV)},
- year = 2018,
- month = sep,
- url = {https://www.merl.com/publications/TR2018-139}
- }
Kocanaogullari, A., Ataer-Cansizoglu, E., "Active Descriptor Learning for Feature Matching", International Workshop on Compact and Efficient Feature Representation and Learning in Computer Vision, September 2018.
BibTeX TR2018-132 PDF
- @inproceedings{Kocanaogullari2018sep,
- author = {Kocanaogullari, Aziz and Ataer-Cansizoglu, Esra},
- title = {{Active Descriptor Learning for Feature Matching}},
- booktitle = {International Workshop on Compact and Efficient Feature Representation and Learning in Computer Vision},
- year = 2018,
- month = sep,
- url = {https://www.merl.com/publications/TR2018-132}
- }
Jones, M.J., Broad, A., Lee, T.-Y., "Recurrent Multi-frame Single Shot Detector for Video Object Detection", British Machine Vision Conference (BMVC), September 2018.
BibTeX TR2018-137 PDF
- @inproceedings{Jones2018sep,
- author = {Jones, Michael J. and Broad, Alexander and Lee, Teng-Yok},
- title = {{Recurrent Multi-frame Single Shot Detector for Video Object Detection}},
- booktitle = {British Machine Vision Conference (BMVC)},
- year = 2018,
- month = sep,
- url = {https://www.merl.com/publications/TR2018-137}
- }
Wang, Z.-Q., Le Roux, J., Wang, D., Hershey, J., "End-to-End Speech Separation with Unfolded Iterative Phase Reconstruction", Interspeech, September 2018.
BibTeX TR2018-135 PDF
- @inproceedings{Wang2018sep,
- author = {Wang, Zhong-Qiu and {Le Roux}, Jonathan and Wang, DeLiang and Hershey, John},
- title = {{End-to-End Speech Separation with Unfolded Iterative Phase Reconstruction}},
- booktitle = {Interspeech},
- year = 2018,
- month = sep,
- url = {https://www.merl.com/publications/TR2018-135}
- }
Watanabe, S., Hori, T., Karita, S., Hayashi, T., Nishitoba, J., Unno, Y., Enrique Yalta Soplin, N., Heymann, J., Wiesner, M., Chen, N., Renduchintala, A., Ochiai, T., "ESPnet: End-to-End Speech Processing Toolkit", Interspeech, September 2018.
BibTeX TR2018-136 PDF
- @inproceedings{Watanabe2018sep,
- author = {Watanabe, Shinji and Hori, Takaaki and Karita, Shigeki and Hayashi, Tomoki and Nishitoba, Jiro and Unno, Yuya and Enrique Yalta Soplin, Nelson and Heymann, Jahn and Wiesner, Matthew and Chen, Nanxin and Renduchintala, Adithya and Ochiai, Tsubasa},
- title = {{ESPnet: End-to-End Speech Processing Toolkit}},
- booktitle = {Interspeech},
- year = 2018,
- month = sep,
- url = {https://www.merl.com/publications/TR2018-136}
- }
Ataer-Cansizoglu, E., Jones, M.J., Zhang, Z., Sullivan, A., "Verification of Very Low-Resolution Faces Using An Identity-Preserving Deep Face Super-resolution Network", arXiv, August 2018.
BibTeX arXiv
- @article{Ataer-Cansizoglu2018aug,
- author = {Ataer-Cansizoglu, Esra and Jones, Michael J. and Zhang, Ziming and Sullivan, Alan},
- title = {{Verification of Very Low-Resolution Faces Using An Identity-Preserving Deep Face Super-resolution Network}},
- journal = {arXiv},
- year = 2018,
- month = aug,
- url = {https://arxiv.org/abs/1903.10974}
- }
Zhang, Z., "LMKL-Net: A Fast Localized Multiple Kernel Learning Solver via Deep Neural Networks", arXiv, July 12, 2018.
BibTeX arXiv
- @article{Zhang2018jul,
- author = {Zhang, Ziming},
- title = {{LMKL-Net: A Fast Localized Multiple Kernel Learning Solver via Deep Neural Networks}},
- journal = {arXiv},
- year = 2018,
- month = jul,
- url = {https://arxiv.org/abs/1805.08656}
- }
Pan, Y., Farahmand, A.-M., White, M., Nabi, S., Grover, P., Nikovski, D.N., "Reinforcement Learning with Function-Valued Action Spaces for Partial Differential Equation Control", International Conference on Machine Learning (ICML), July 2018.
BibTeX TR2018-101 PDF
- @inproceedings{Pan2018jul,
- author = {Pan, Yangchen and Farahmand, Amir-massoud and White, Martha and Nabi, Saleh and Grover, Piyush and Nikovski, Daniel N.},
- title = {{Reinforcement Learning with Function-Valued Action Spaces for Partial Differential Equation Control}},
- booktitle = {International Conference on Machine Learning (ICML)},
- year = 2018,
- month = jul,
- url = {https://www.merl.com/publications/TR2018-101}
- }
Barker, J., Marxer, R., Vincent, E., Watanabe, S., "The CHiME challenges: Robust speech recognition in everyday environments" in New Era for Robust Speech Recognition: Exploiting Deep Learning, Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R., Eds., chapter 14, Springer, July 2018.
BibTeX
- @incollection{Barker2018jul,
- author = {Barker, Jon and Marxer, Ricard and Vincent, Emmanuel and Watanabe, Shinji},
- title = {{The CHiME challenges: Robust speech recognition in everyday environments}},
- booktitle = {New Era for Robust Speech Recognition: Exploiting Deep Learning},
- year = 2018,
- editor = {Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R.},
- chapter = 14,
- month = jul,
- publisher = {Springer}
- }
Erdogan, H., Hershey, J., Watanabe, S., Le Roux, J., "Deep recurrent networks for separation and recognition of single-channel speech in non-stationary background audio" in New Era for Robust Speech Recognition: Exploiting Deep Learning, Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R., Eds., chapter 7, Springer, July 2018.
BibTeX
- @incollection{Erdogan2018jul,
- author = {Erdogan, Hakan and Hershey, John and Watanabe, Shinji and {Le Roux}, Jonathan},
- title = {{Deep recurrent networks for separation and recognition of single-channel speech in non-stationary background audio}},
- booktitle = {New Era for Robust Speech Recognition: Exploiting Deep Learning},
- year = 2018,
- editor = {Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R.},
- chapter = 7,
- month = jul,
- publisher = {Springer},
- isbn = {978-3-319-64680-0}
- }
Hershey, J., Le Roux, J., Watanabe, S., Wisdom, S., Chen, Z., Isik, Y., "Novel deep architectures in speech processing" in New Era for Robust Speech Recognition: Exploiting Deep Learning, Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R., Eds., chapter 6, Springer, July 9, 2018.
BibTeX
- @incollection{Hershey2018jul,
- author = {Hershey, John and {Le Roux}, Jonathan and Watanabe, Shinji and Wisdom, Scott and Chen, Zhuo and Isik, Yusuf},
- title = {{Novel deep architectures in speech processing}},
- booktitle = {New Era for Robust Speech Recognition: Exploiting Deep Learning},
- year = 2018,
- editor = {Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R.},
- chapter = 6,
- month = jul,
- publisher = {Springer}
- }
Karafiat, M., Vesely, K., Zmolikova, K., Delcroix, M., Watanabe, S., Burget, L., Cernocky, J., Szoke, I., Novotny, O., "Training data augmentation and data selectio" in New Era for Robust Speech Recognition: Exploiting Deep Learning, Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R., Eds., chapter 10, Springer, July 9, 2018.
BibTeX
- @incollection{Karafiat2018jul,
- author = {Karafiat, Martin and Vesely, Karel and Zmolikova, Katerina and Delcroix, Marc and Watanabe, Shinji and Burget, Lukas and Cernocky, Jan and Szoke, Igor and Novotny, Ondrej},
- title = {{Training data augmentation and data selectio}},
- booktitle = {New Era for Robust Speech Recognition: Exploiting Deep Learning},
- year = 2018,
- editor = {Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R.},
- chapter = 10,
- month = jul,
- publisher = {Springer}
- }
Watanabe, S., Hori, T., Miao, Y., Delcroix, M., Metze, F., Hershey, J., "Toolkits for robust speech processing" in New Era for Robust Speech Recognition: Exploiting Deep Learning, Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R., Eds., chapter 14, Springer, July 9, 2018.
BibTeX
- @incollection{Watanabe2018jul,
- author = {Watanabe, Shinji and Hori, Takaaki and Miao, Yajie and Delcroix, Marc and Metze, Florian and Hershey, John},
- title = {{Toolkits for robust speech processing}},
- booktitle = {New Era for Robust Speech Recognition: Exploiting Deep Learning},
- year = 2018,
- editor = {Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R.},
- chapter = 14,
- month = jul,
- publisher = {Springer}
- }
Xiao, X., Watanabe, S., Erdogan, H., Mandel, M., Lu, L., Hershey, J., Seltzer, M., Chen, G., Zhang, Y., Yu, D., "Discriminative beamforming with phase aware neural networks for speech enhancement and recognition" in New Era for Robust Speech Recognition: Exploiting Deep Learning, Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R., Eds., chapter 4, Springer, July 9, 2018.
BibTeX
- @incollection{Xiao2018jul2,
- author = {Xiao, Xiong and Watanabe, Shinji and Erdogan, Hakan and Mandel, Michael and Lu, Liang and Hershey, John and Seltzer, Mike and Chen, Guoguo and Zhang, Yu and Yu, Dong},
- title = {{Discriminative beamforming with phase aware neural networks for speech enhancement and recognition}},
- booktitle = {New Era for Robust Speech Recognition: Exploiting Deep Learning},
- year = 2018,
- editor = {Watanabe, S. and Delcroix, M. and Metze, F. and Hershey, J.R.},
- chapter = 4,
- month = jul,
- publisher = {Springer}
- }
Wang, P., Fu, H., Koike-Akino, T., Orlik, P.V., "Multi-Layer Terahertz Imaging of Non-Overlapping Contents", IEEE Sensor Array and Multi-Channel Signal Processing Workshop (IEEE SAM), DOI: 10.1109/SAM.2018.8448438, July 2018, pp. 652-656.
BibTeX TR2018-098 PDF
- @inproceedings{Wang2018jul2,
- author = {Wang, Pu and Fu, Haoyu and Koike-Akino, Toshiaki and Orlik, Philip V.},
- title = {{Multi-Layer Terahertz Imaging of Non-Overlapping Contents}},
- booktitle = {IEEE Sensor Array and Multi-Channel Signal Processing Workshop (IEEE SAM)},
- year = 2018,
- pages = {652--656},
- month = jul,
- doi = {10.1109/SAM.2018.8448438},
- url = {https://www.merl.com/publications/TR2018-098}
- }
Liu, J., Guo, J., Orlik, P.V., Shibata, M., Nakahara, D., Mii, S., Takac, M., "Anomaly Detection in Manufacturing Systems Using Structured Neural Networks", IEEE World Congress on Intelligent Control and Automation, DOI: 10.1109/WCICA.2018.8630692, July 2018, pp. 175-180.
BibTeX TR2018-097 PDF
- @inproceedings{Liu2018jul2,
- author = {Liu, Jie and Guo, Jianlin and Orlik, Philip V. and Shibata, Masahiko and Nakahara, Daiki and Mii, Satoshi and Takac, Martin},
- title = {{Anomaly Detection in Manufacturing Systems Using Structured Neural Networks}},
- booktitle = {IEEE World Congress on Intelligent Control and Automation},
- year = 2018,
- pages = {175--180},
- month = jul,
- doi = {10.1109/WCICA.2018.8630692},
- url = {https://www.merl.com/publications/TR2018-097}
- }
Koike-Akino, T., Millar, D.S., Parsons, K., Kojima, K., "Fiber Nonlinearity Equalization with Multi-Label Deep Learning Scalable to High-Order DP-QAM", Signal Processing in Photonic Communications (SPPCom), DOI: 10.1364/SPPCOM.2018.SpM4G.1, July 2018.
BibTeX TR2018-047 PDF
- @inproceedings{Koike-Akino2018jul3,
- author = {Koike-Akino, Toshiaki and Millar, David S. and Parsons, Kieran and Kojima, Keisuke},
- title = {{Fiber Nonlinearity Equalization with Multi-Label Deep Learning Scalable to High-Order DP-QAM}},
- booktitle = {Signal Processing in Photonic Communications (SPPCom)},
- year = 2018,
- month = jul,
- doi = {10.1364/SPPCOM.2018.SpM4G.1},
- url = {https://www.merl.com/publications/TR2018-047}
- }
Shen, Y., Feng, C., Yang, Y., Tian, D., "Mining Point Cloud Local Structures by Kernel Correlation and Graph Pooling", IEEE Conference on Computer Vision and Pattern Recognition (CVPR), June 2018.
BibTeX TR2018-041 PDF Software
- @inproceedings{Shen2018jun,
- author = {Shen, Yiru and Feng, Chen and Yang, Yaoqing and Tian, Dong},
- title = {{Mining Point Cloud Local Structures by Kernel Correlation and Graph Pooling}},
- booktitle = {IEEE Conference on Computer Vision and Pattern Recognition (CVPR)},
- year = 2018,
- month = jun,
- url = {https://www.merl.com/publications/TR2018-041}
- }