3D ShapeNets: A deep representation for volumetric shapes

3D ShapeNets: A deep representation for volumetric shapes. Wu, Z., Song, S., Khosla, A., Yu, F., Zhang, L., Tang, X., & Xiao, J. Proceedings of the IEEE Computer Society Conference on Computer Vision and Pattern Recognition, 07-12-June:1912-1920, 2015.

Paper doi abstract bibtex

3D shape is a crucial but heavily underutilized cue in today's computer vision systems, mostly due to the lack of a good generic shape representation. With the recent availability of inexpensive 2.5D depth sensors (e.g. Microsoft Kinect), it is becoming increasingly important to have a powerful 3D shape representation in the loop. Apart from category recognition, recovering full 3D shapes from view-based 2.5D depth maps is also a critical part of visual understanding. To this end, we propose to represent a geometric 3D shape as a probability distribution of binary variables on a 3D voxel grid, using a Convolutional Deep Belief Network. Our model, 3D ShapeNets, learns the distribution of complex 3D shapes across different object categories and arbitrary poses from raw CAD data, and discovers hierarchical compositional part representation automatically. It naturally supports joint object recognition and shape completion from 2.5D depth maps, and it enables active object recognition through view planning. To train our 3D deep learning model, we construct ModelNet - a large-scale 3D CAD model dataset. Extensive experiments show that our 3D deep representation enables significant performance improvement over the-state-of-the-arts in a variety of tasks.

@article{
 title = {3D ShapeNets: A deep representation for volumetric shapes},
 type = {article},
 year = {2015},
 pages = {1912-1920},
 volume = {07-12-June},
 id = {a0d1d401-c93f-34c8-b08d-11b3626e7c61},
 created = {2021-07-26T12:19:39.684Z},
 file_attached = {true},
 profile_id = {ad172e55-c0e8-3aa4-8465-09fac4d5f5c8},
 group_id = {1ff583c0-be37-34fa-9c04-73c69437d354},
 last_modified = {2021-07-26T12:19:45.070Z},
 read = {false},
 starred = {false},
 authored = {false},
 confirmed = {true},
 hidden = {false},
 citation_key = {Wu2015},
 folder_uuids = {990ab628-0917-4e89-b071-24bf1f44fad6,4f36a0a5-b08a-4f70-b020-4daf83cb0507},
 private_publication = {false},
 abstract = {3D shape is a crucial but heavily underutilized cue in today's computer vision systems, mostly due to the lack of a good generic shape representation. With the recent availability of inexpensive 2.5D depth sensors (e.g. Microsoft Kinect), it is becoming increasingly important to have a powerful 3D shape representation in the loop. Apart from category recognition, recovering full 3D shapes from view-based 2.5D depth maps is also a critical part of visual understanding. To this end, we propose to represent a geometric 3D shape as a probability distribution of binary variables on a 3D voxel grid, using a Convolutional Deep Belief Network. Our model, 3D ShapeNets, learns the distribution of complex 3D shapes across different object categories and arbitrary poses from raw CAD data, and discovers hierarchical compositional part representation automatically. It naturally supports joint object recognition and shape completion from 2.5D depth maps, and it enables active object recognition through view planning. To train our 3D deep learning model, we construct ModelNet - a large-scale 3D CAD model dataset. Extensive experiments show that our 3D deep representation enables significant performance improvement over the-state-of-the-arts in a variety of tasks.},
 bibtype = {article},
 author = {Wu, Zhirong and Song, Shuran and Khosla, Aditya and Yu, Fisher and Zhang, Linguang and Tang, Xiaoou and Xiao, Jianxiong},
 doi = {10.1109/CVPR.2015.7298801},
 journal = {Proceedings of the IEEE Computer Society Conference on Computer Vision and Pattern Recognition}
}

Downloads: 0

{"_id":"KEnbhL67thuhcSNJX","bibbaseid":"wu-song-khosla-yu-zhang-tang-xiao-3dshapenetsadeeprepresentationforvolumetricshapes-2015","author_short":["Wu, Z.","Song, S.","Khosla, A.","Yu, F.","Zhang, L.","Tang, X.","Xiao, J."],"bibdata":{"title":"3D ShapeNets: A deep representation for volumetric shapes","type":"article","year":"2015","pages":"1912-1920","volume":"07-12-June","id":"a0d1d401-c93f-34c8-b08d-11b3626e7c61","created":"2021-07-26T12:19:39.684Z","file_attached":"true","profile_id":"ad172e55-c0e8-3aa4-8465-09fac4d5f5c8","group_id":"1ff583c0-be37-34fa-9c04-73c69437d354","last_modified":"2021-07-26T12:19:45.070Z","read":false,"starred":false,"authored":false,"confirmed":"true","hidden":false,"citation_key":"Wu2015","folder_uuids":"990ab628-0917-4e89-b071-24bf1f44fad6,4f36a0a5-b08a-4f70-b020-4daf83cb0507","private_publication":false,"abstract":"3D shape is a crucial but heavily underutilized cue in today's computer vision systems, mostly due to the lack of a good generic shape representation. With the recent availability of inexpensive 2.5D depth sensors (e.g. Microsoft Kinect), it is becoming increasingly important to have a powerful 3D shape representation in the loop. Apart from category recognition, recovering full 3D shapes from view-based 2.5D depth maps is also a critical part of visual understanding. To this end, we propose to represent a geometric 3D shape as a probability distribution of binary variables on a 3D voxel grid, using a Convolutional Deep Belief Network. Our model, 3D ShapeNets, learns the distribution of complex 3D shapes across different object categories and arbitrary poses from raw CAD data, and discovers hierarchical compositional part representation automatically. It naturally supports joint object recognition and shape completion from 2.5D depth maps, and it enables active object recognition through view planning. To train our 3D deep learning model, we construct ModelNet - a large-scale 3D CAD model dataset. Extensive experiments show that our 3D deep representation enables significant performance improvement over the-state-of-the-arts in a variety of tasks.","bibtype":"article","author":"Wu, Zhirong and Song, Shuran and Khosla, Aditya and Yu, Fisher and Zhang, Linguang and Tang, Xiaoou and Xiao, Jianxiong","doi":"10.1109/CVPR.2015.7298801","journal":"Proceedings of the IEEE Computer Society Conference on Computer Vision and Pattern Recognition","bibtex":"@article{\n title = {3D ShapeNets: A deep representation for volumetric shapes},\n type = {article},\n year = {2015},\n pages = {1912-1920},\n volume = {07-12-June},\n id = {a0d1d401-c93f-34c8-b08d-11b3626e7c61},\n created = {2021-07-26T12:19:39.684Z},\n file_attached = {true},\n profile_id = {ad172e55-c0e8-3aa4-8465-09fac4d5f5c8},\n group_id = {1ff583c0-be37-34fa-9c04-73c69437d354},\n last_modified = {2021-07-26T12:19:45.070Z},\n read = {false},\n starred = {false},\n authored = {false},\n confirmed = {true},\n hidden = {false},\n citation_key = {Wu2015},\n folder_uuids = {990ab628-0917-4e89-b071-24bf1f44fad6,4f36a0a5-b08a-4f70-b020-4daf83cb0507},\n private_publication = {false},\n abstract = {3D shape is a crucial but heavily underutilized cue in today's computer vision systems, mostly due to the lack of a good generic shape representation. With the recent availability of inexpensive 2.5D depth sensors (e.g. Microsoft Kinect), it is becoming increasingly important to have a powerful 3D shape representation in the loop. Apart from category recognition, recovering full 3D shapes from view-based 2.5D depth maps is also a critical part of visual understanding. To this end, we propose to represent a geometric 3D shape as a probability distribution of binary variables on a 3D voxel grid, using a Convolutional Deep Belief Network. Our model, 3D ShapeNets, learns the distribution of complex 3D shapes across different object categories and arbitrary poses from raw CAD data, and discovers hierarchical compositional part representation automatically. It naturally supports joint object recognition and shape completion from 2.5D depth maps, and it enables active object recognition through view planning. To train our 3D deep learning model, we construct ModelNet - a large-scale 3D CAD model dataset. Extensive experiments show that our 3D deep representation enables significant performance improvement over the-state-of-the-arts in a variety of tasks.},\n bibtype = {article},\n author = {Wu, Zhirong and Song, Shuran and Khosla, Aditya and Yu, Fisher and Zhang, Linguang and Tang, Xiaoou and Xiao, Jianxiong},\n doi = {10.1109/CVPR.2015.7298801},\n journal = {Proceedings of the IEEE Computer Society Conference on Computer Vision and Pattern Recognition}\n}","author_short":["Wu, Z.","Song, S.","Khosla, A.","Yu, F.","Zhang, L.","Tang, X.","Xiao, J."],"urls":{"Paper":"https://bibbase.org/service/mendeley/bfbbf840-4c42-3914-a463-19024f50b30c/file/03c8ca53-802e-7efd-6b2e-3360a1638e5e/3D_ShapeNets_A_Deep_Representation_for_Volumetric_Shape_Modeling.pdf.pdf"},"biburl":"https://bibbase.org/service/mendeley/bfbbf840-4c42-3914-a463-19024f50b30c","bibbaseid":"wu-song-khosla-yu-zhang-tang-xiao-3dshapenetsadeeprepresentationforvolumetricshapes-2015","role":"author","metadata":{"authorlinks":{}}},"bibtype":"article","biburl":"https://bibbase.org/service/mendeley/bfbbf840-4c42-3914-a463-19024f50b30c","dataSources":["Qsjejxqv2Xv9k84sx","ya2CyA73rpZseyrZ8","2252seNhipfTmjEBQ"],"keywords":[],"search_terms":["shapenets","deep","representation","volumetric","shapes","wu","song","khosla","yu","zhang","tang","xiao"],"title":"3D ShapeNets: A deep representation for volumetric shapes","year":2015}