Anurag1734/cuda-error-resolution-analysis
07
1[2 {3 "post_stream": {4 "posts": [5 {6 "id": 154159,7 "name": "",8 "username": "Shubhankar",9 "avatar_template": "/letter_avatar_proxy/v4/letter/s/b9bd4f/{size}.png",10 "created_at": "2019-12-20T20:39:56.176Z",11 "cooked": "<p>How to use NVIDIA AUTO TUNE with pytorch.</p>",12 "post_number": 1,13 "post_type": 1,14 "posts_count": 4,15 "updated_at": "2019-12-20T20:39:56.176Z",16 "reply_count": 0,17 "reply_to_post_number": null,18 "quote_count": 0,19 "incoming_link_count": 745,20 "reads": 26,21 "readers_count": 25,22 "score": 3725.2,23 "yours": false,24 "topic_id": 64669,25 "topic_slug": "how-to-use-auto-tune-with-pytorch",26 "display_username": "",27 "primary_group_name": null,28 "flair_name": null,29 "flair_url": null,30 "flair_bg_color": null,31 "flair_color": null,32 "flair_group_id": null,33 "badges_granted": [],34 "version": 1,35 "can_edit": false,36 "can_delete": false,37 "can_recover": false,38 "can_see_hidden_post": false,39 "can_wiki": false,40 "read": true,41 "user_title": "",42 "bookmarked": false,43 "actions_summary": [],44 "moderator": false,45 "admin": false,46 "staff": false,47 "user_id": 24042,48 "hidden": false,49 "trust_level": 1,50 "deleted_at": null,51 "user_deleted": false,52 "edit_reason": null,53 "can_view_edit_history": true,54 "wiki": false,55 "post_url": "/t/how-to-use-auto-tune-with-pytorch/64669/1",56 "can_accept_answer": false,57 "can_unaccept_answer": false,58 "accepted_answer": false,59 "topic_accepted_answer": true,60 "can_vote": false61 },62 {63 "id": 154388,64 "name": "Alban D",65 "username": "albanD",66 "avatar_template": "/user_avatar/discuss.pytorch.org/alband/{size}/215_2.png",67 "created_at": "2019-12-22T13:47:57.180Z",68 "cooked": "<p>Hi,</p>\n<p>I am not familiar with this? Could you explain what it is and what it does please?</p>",69 "post_number": 2,70 "post_type": 1,71 "posts_count": 4,72 "updated_at": "2019-12-22T13:47:57.180Z",73 "reply_count": 0,74 "reply_to_post_number": null,75 "quote_count": 0,76 "incoming_link_count": 5,77 "reads": 24,78 "readers_count": 23,79 "score": 29.8,80 "yours": false,81 "topic_id": 64669,82 "topic_slug": "how-to-use-auto-tune-with-pytorch",83 "display_username": "Alban D",84 "primary_group_name": null,85 "flair_name": null,86 "flair_url": null,87 "flair_bg_color": null,88 "flair_color": null,89 "flair_group_id": null,90 "badges_granted": [],91 "version": 1,92 "can_edit": false,93 "can_delete": false,94 "can_recover": false,95 "can_see_hidden_post": false,96 "can_wiki": false,97 "read": true,98 "user_title": "",99 "bookmarked": false,100 "actions_summary": [],101 "moderator": true,102 "admin": true,103 "staff": true,104 "user_id": 211,105 "hidden": false,106 "trust_level": 4,107 "deleted_at": null,108 "user_deleted": false,109 "edit_reason": null,110 "can_view_edit_history": true,111 "wiki": false,112 "post_url": "/t/how-to-use-auto-tune-with-pytorch/64669/2",113 "can_accept_answer": false,114 "can_unaccept_answer": false,115 "accepted_answer": false,116 "topic_accepted_answer": true117 },118 {119 "id": 154407,120 "name": "",121 "username": "Shubhankar",122 "avatar_template": "/letter_avatar_proxy/v4/letter/s/b9bd4f/{size}.png",123 "created_at": "2019-12-22T17:23:02.009Z",124 "cooked": "<aside class=\"onebox githubissue\">\n <header class=\"source\">\n <a href=\"https://github.com/tensorflow/tensorflow/issues/12871#issuecomment-370059148\" target=\"_blank\" rel=\"nofollow noopener\">github.com/tensorflow/tensorflow</a>\n </header>\n <article class=\"onebox-body\">\n <div class=\"github-row\">\n <div class=\"github-icon-container\" title=\"Issue\">\n\t <svg width=\"60\" height=\"60\" class=\"github-icon\" viewbox=\"0 0 14 16\" aria-hidden=\"true\"><path d=\"M7 2.3c3.14 0 5.7 2.56 5.7 5.7s-2.56 5.7-5.7 5.7A5.71 5.71 0 0 1 1.3 8c0-3.14 2.56-5.7 5.7-5.7zM7 1C3.14 1 0 4.14 0 8s3.14 7 7 7 7-3.14 7-7-3.14-7-7-7zm1 3H6v5h2V4zm0 6H6v2h2v-2z\"></path></svg>\n </div>\n\n <div class=\"github-info-container\">\n <h4>\n <a href=\"https://github.com/tensorflow/tensorflow/issues/12871#issuecomment-370059148\" target=\"_blank\" rel=\"nofollow noopener\">About Deterministic Behaviour of GPU implementation of tensorflow</a>\n </h4>\n\n <div class=\"github-info\">\n <div class=\"date\">\n opened <span class=\"discourse-local-date\" data-format=\"ll\" data-date=\"2017-09-07\" data-time=\"08:07:09\" data-timezone=\"UTC\">08:07AM - 07 Sep 17 UTC</span>\n </div>\n\n <div class=\"date\">\n closed <span class=\"discourse-local-date\" data-format=\"ll\" data-date=\"2018-09-21\" data-time=\"18:57:17\" data-timezone=\"UTC\">06:57PM - 21 Sep 18 UTC</span>\n </div>\n\n <div class=\"user\">\n <a href=\"https://github.com/antares1987\" target=\"_blank\" rel=\"nofollow noopener\">\n <img alt=\"antares1987\" src=\"https://avatars0.githubusercontent.com/u/31724739?v=4\" class=\"onebox-avatar-inline\" width=\"20\" height=\"20\">\n antares1987\n </a>\n </div>\n </div>\n </div>\n</div>\n\n<div class=\"github-row\">\n <p class=\"github-content\">OS Platform and Distribution (e.g., Linux Ubuntu 16.04): Debian\nTensorFlow version (use command below):('v1.3.0-rc2-20-g0787eee', '1.3.0')\nCUDA/cuDNN version: 8.0\nGPU model and memory: GeForce GTX...</p>\n</div>\n\n<div class=\"labels\">\n <span style=\"display:inline-block;margin-top:2px;background-color: #B8B8B8;padding: 2px;border-radius: 4px;color: #fff;margin-left: 3px;\">stat:awaiting response</span>\n</div>\n\n </article>\n <div class=\"onebox-metadata\">\n \n \n </div>\n <div style=\"clear: both\"></div>\n</aside>\n",125 "post_number": 3,126 "post_type": 1,127 "posts_count": 4,128 "updated_at": "2019-12-22T17:23:02.009Z",129 "reply_count": 1,130 "reply_to_post_number": null,131 "quote_count": 0,132 "incoming_link_count": 5,133 "reads": 23,134 "readers_count": 22,135 "score": 34.6,136 "yours": false,137 "topic_id": 64669,138 "topic_slug": "how-to-use-auto-tune-with-pytorch",139 "display_username": "",140 "primary_group_name": null,141 "flair_name": null,142 "flair_url": null,143 "flair_bg_color": null,144 "flair_color": null,145 "flair_group_id": null,146 "badges_granted": [],147 "version": 1,148 "can_edit": false,149 "can_delete": false,150 "can_recover": false,151 "can_see_hidden_post": false,152 "can_wiki": false,153 "link_counts": [154 {155 "url": "https://github.com/tensorflow/tensorflow/issues/12871#issuecomment-370059148",156 "internal": false,157 "reflection": false,158 "title": "About Deterministic Behaviour of GPU implementation of tensorflow · Issue #12871 · tensorflow/tensorflow · GitHub",159 "clicks": 10160 },161 {162 "url": "https://github.com/antares1987",163 "internal": false,164 "reflection": false,165 "title": "antares1987 · GitHub",166 "clicks": 0167 }168 ],169 "read": true,170 "user_title": "",171 "bookmarked": false,172 "actions_summary": [],173 "moderator": false,174 "admin": false,175 "staff": false,176 "user_id": 24042,177 "hidden": false,178 "trust_level": 1,179 "deleted_at": null,180 "user_deleted": false,181 "edit_reason": null,182 "can_view_edit_history": true,183 "wiki": false,184 "post_url": "/t/how-to-use-auto-tune-with-pytorch/64669/3",185 "can_accept_answer": false,186 "can_unaccept_answer": false,187 "accepted_answer": false,188 "topic_accepted_answer": true189 },190 {191 "id": 154475,192 "name": "",193 "username": "ptrblck",194 "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png",195 "created_at": "2019-12-23T07:04:18.869Z",196 "cooked": "<p>I’m not familiar with TensorFlow, but it seems their “auto tune” functionality corresponds to <code>torch.backends.cudnn.benchmark = True</code> mode.</p>",197 "post_number": 4,198 "post_type": 1,199 "posts_count": 4,200 "updated_at": "2020-01-06T05:18:09.159Z",201 "reply_count": 0,202 "reply_to_post_number": 3,203 "quote_count": 0,204 "incoming_link_count": 7,205 "reads": 19,206 "readers_count": 18,207 "score": 83.8,208 "yours": false,209 "topic_id": 64669,210 "topic_slug": "how-to-use-auto-tune-with-pytorch",211 "display_username": "",212 "primary_group_name": null,213 "flair_name": null,214 "flair_url": null,215 "flair_bg_color": null,216 "flair_color": null,217 "flair_group_id": null,218 "badges_granted": [],219 "version": 1,220 "can_edit": false,221 "can_delete": false,222 "can_recover": false,223 "can_see_hidden_post": false,224 "can_wiki": false,225 "read": true,226 "user_title": "",227 "reply_to_user": {228 "id": 24042,229 "username": "Shubhankar",230 "name": "",231 "avatar_template": "/letter_avatar_proxy/v4/letter/s/b9bd4f/{size}.png"232 },233 "bookmarked": false,234 "actions_summary": [235 {236 "id": 2,237 "count": 1238 }239 ],240 "moderator": true,241 "admin": true,242 "staff": true,243 "user_id": 3534,244 "hidden": false,245 "trust_level": 2,246 "deleted_at": null,247 "user_deleted": false,248 "edit_reason": null,249 "can_view_edit_history": true,250 "wiki": false,251 "post_url": "/t/how-to-use-auto-tune-with-pytorch/64669/4",252 "can_accept_answer": false,253 "can_unaccept_answer": false,254 "accepted_answer": true,255 "topic_accepted_answer": true256 }257 ],258 "stream": [259 154159,260 154388,261 154407,262 154475263 ]264 },265 "timeline_lookup": [266 [267 1,268 2136269 ],270 [271 2,272 2134273 ]274 ],275 "suggested_topics": [276 {277 "fancy_title": "ExecuTorch: Variable-Length Inputs for Export and F.maxpool1d",278 "id": 218652,279 "title": "ExecuTorch: Variable-Length Inputs for Export and F.maxpool1d",280 "slug": "executorch-variable-length-inputs-for-export-and-f-maxpool1d",281 "posts_count": 1,282 "reply_count": 0,283 "highest_post_number": 1,284 "image_url": null,285 "created_at": "2025-04-06T02:00:56.155Z",286 "last_posted_at": "2025-04-06T02:00:56.195Z",287 "bumped": true,288 "bumped_at": "2025-04-06T02:00:56.195Z",289 "archetype": "regular",290 "unseen": false,291 "pinned": false,292 "unpinned": null,293 "visible": true,294 "closed": false,295 "archived": false,296 "bookmarked": null,297 "liked": null,298 "tags_descriptions": {},299 "like_count": 0,300 "views": 19,301 "category_id": 1,302 "featured_link": null,303 "has_accepted_answer": false,304 "posters": [305 {306 "extras": "latest single",307 "description": "Original Poster, Most Recent Poster",308 "user": {309 "id": 83653,310 "username": "cba913",311 "name": "",312 "avatar_template": "/user_avatar/discuss.pytorch.org/cba913/{size}/76501_2.png",313 "trust_level": 0314 }315 }316 ]317 },318 {319 "fancy_title": "How to approach a real-life problem while using rainfall data",320 "id": 212728,321 "title": "How to approach a real-life problem while using rainfall data",322 "slug": "how-to-approach-a-real-life-problem-while-using-rainfall-data",323 "posts_count": 1,324 "reply_count": 0,325 "highest_post_number": 1,326 "image_url": null,327 "created_at": "2024-11-09T07:21:57.841Z",328 "last_posted_at": "2024-11-09T07:21:57.888Z",329 "bumped": true,330 "bumped_at": "2024-11-09T07:21:57.888Z",331 "archetype": "regular",332 "unseen": false,333 "pinned": false,334 "unpinned": null,335 "visible": true,336 "closed": false,337 "archived": false,338 "bookmarked": null,339 "liked": null,340 "tags_descriptions": {},341 "like_count": 0,342 "views": 25,343 "category_id": 1,344 "featured_link": null,345 "has_accepted_answer": false,346 "posters": [347 {348 "extras": "latest single",349 "description": "Original Poster, Most Recent Poster",350 "user": {351 "id": 80781,352 "username": "Ritam_Pradhan",353 "name": "Ritam Pradhan",354 "avatar_template": "/user_avatar/discuss.pytorch.org/ritam_pradhan/{size}/73876_2.png",355 "trust_level": 1356 }357 }358 ]359 },360 {361 "fancy_title": "CUDA out of memory while using Llama3.1-8B for inference",362 "id": 213808,363 "title": "CUDA out of memory while using Llama3.1-8B for inference",364 "slug": "cuda-out-of-memory-while-using-llama3-1-8b-for-inference",365 "posts_count": 1,366 "reply_count": 0,367 "highest_post_number": 1,368 "image_url": null,369 "created_at": "2024-12-04T17:52:39.019Z",370 "last_posted_at": "2024-12-04T17:52:39.070Z",371 "bumped": true,372 "bumped_at": "2024-12-04T17:52:39.070Z",373 "archetype": "regular",374 "unseen": false,375 "pinned": false,376 "unpinned": null,377 "visible": true,378 "closed": false,379 "archived": false,380 "bookmarked": null,381 "liked": null,382 "tags_descriptions": {},383 "like_count": 0,384 "views": 96,385 "category_id": 1,386 "featured_link": null,387 "has_accepted_answer": false,388 "posters": [389 {390 "extras": "latest single",391 "description": "Original Poster, Most Recent Poster",392 "user": {393 "id": 70667,394 "username": "tomwagstaff-opml",395 "name": "Tom Wagstaff",396 "avatar_template": "/user_avatar/discuss.pytorch.org/tomwagstaff-opml/{size}/65211_2.png",397 "trust_level": 0398 }399 }400 ]401 },402 {403 "fancy_title": "ROCm: hipBLASLt error with gfx1103",404 "id": 212721,405 "title": "ROCm: hipBLASLt error with gfx1103",406 "slug": "rocm-hipblaslt-error-with-gfx1103",407 "posts_count": 2,408 "reply_count": 0,409 "highest_post_number": 2,410 "image_url": null,411 "created_at": "2024-11-09T01:43:15.446Z",412 "last_posted_at": "2025-03-03T19:42:45.282Z",413 "bumped": true,414 "bumped_at": "2025-03-03T19:42:45.282Z",415 "archetype": "regular",416 "unseen": false,417 "pinned": false,418 "unpinned": null,419 "visible": true,420 "closed": false,421 "archived": false,422 "bookmarked": null,423 "liked": null,424 "tags_descriptions": {},425 "like_count": 0,426 "views": 1070,427 "category_id": 1,428 "featured_link": null,429 "has_accepted_answer": false,430 "posters": [431 {432 "extras": null,433 "description": "Original Poster",434 "user": {435 "id": 78358,436 "username": "lumie",437 "name": "",438 "avatar_template": "/letter_avatar_proxy/v4/letter/l/ea5d25/{size}.png",439 "trust_level": 1440 }441 },442 {443 "extras": "latest",444 "description": "Most Recent Poster",445 "user": {446 "id": 82743,447 "username": "fngarrett",448 "name": "Garrett",449 "avatar_template": "/letter_avatar_proxy/v4/letter/f/e79b87/{size}.png",450 "trust_level": 1451 }452 }453 ]454 },455 {456 "fancy_title": "Tuning a network, subset of data, one non-frozen layer",457 "id": 218788,458 "title": "Tuning a network, subset of data, one non-frozen layer",459 "slug": "tuning-a-network-subset-of-data-one-non-frozen-layer",460 "posts_count": 1,461 "reply_count": 0,462 "highest_post_number": 1,463 "image_url": null,464 "created_at": "2025-04-07T05:54:35.100Z",465 "last_posted_at": "2025-04-07T05:54:35.140Z",466 "bumped": true,467 "bumped_at": "2025-04-07T05:54:35.140Z",468 "archetype": "regular",469 "unseen": false,470 "pinned": false,471 "unpinned": null,472 "visible": true,473 "closed": false,474 "archived": false,475 "bookmarked": null,476 "liked": null,477 "tags_descriptions": {},478 "like_count": 0,479 "views": 18,480 "category_id": 1,481 "featured_link": null,482 "has_accepted_answer": false,483 "posters": [484 {485 "extras": "latest single",486 "description": "Original Poster, Most Recent Poster",487 "user": {488 "id": 78333,489 "username": "emerth",490 "name": "",491 "avatar_template": "/user_avatar/discuss.pytorch.org/emerth/{size}/75379_2.png",492 "trust_level": 1493 }494 }495 ]496 }497 ],498 "tags_descriptions": {},499 "fancy_title": "How to use AUTO_TUNE with pytorch?",500 "id": 64669,501 "title": "How to use AUTO_TUNE with pytorch?",502 "posts_count": 4,503 "created_at": "2019-12-20T20:39:56.126Z",504 "views": 1525,505 "reply_count": 1,506 "like_count": 1,507 "last_posted_at": "2019-12-23T07:04:18.869Z",508 "visible": true,509 "closed": false,510 "archived": false,511 "has_summary": false,512 "archetype": "regular",513 "slug": "how-to-use-auto-tune-with-pytorch",514 "category_id": 1,515 "word_count": 56,516 "deleted_at": null,517 "user_id": 24042,518 "featured_link": null,519 "pinned_globally": false,520 "pinned_at": null,521 "pinned_until": null,522 "image_url": null,523 "slow_mode_seconds": 0,524 "draft": null,525 "draft_key": "topic_64669",526 "draft_sequence": null,527 "unpinned": null,528 "pinned": false,529 "current_post_number": 1,530 "highest_post_number": 4,531 "deleted_by": null,532 "actions_summary": [533 {534 "id": 4,535 "count": 0,536 "hidden": false,537 "can_act": false538 },539 {540 "id": 8,541 "count": 0,542 "hidden": false,543 "can_act": false544 },545 {546 "id": 10,547 "count": 0,548 "hidden": false,549 "can_act": false550 },551 {552 "id": 7,553 "count": 0,554 "hidden": false,555 "can_act": false556 }557 ],558 "chunk_size": 20,559 "bookmarked": false,560 "topic_timer": null,561 "message_bus_last_id": 0,562 "participant_count": 3,563 "show_read_indicator": false,564 "thumbnails": null,565 "slow_mode_enabled_until": null,566 "accepted_answer": {567 "post_number": 4,568 "username": "ptrblck",569 "name": "",570 "excerpt": "I’m not familiar with TensorFlow, but it seems their “auto tune” functionality corresponds to torch.backends.cudnn.benchmark = True mode."571 },572 "can_vote": false,573 "vote_count": 0,574 "user_voted": false,575 "discourse_zendesk_plugin_zendesk_id": null,576 "discourse_zendesk_plugin_zendesk_url": "https://your-url.zendesk.com/agent/tickets/",577 "details": {578 "can_edit": false,579 "notification_level": 1,580 "participants": [581 {582 "id": 24042,583 "username": "Shubhankar",584 "name": "",585 "avatar_template": "/letter_avatar_proxy/v4/letter/s/b9bd4f/{size}.png",586 "post_count": 2,587 "primary_group_name": null,588 "flair_name": null,589 "flair_url": null,590 "flair_color": null,591 "flair_bg_color": null,592 "flair_group_id": null,593 "trust_level": 1594 },595 {596 "id": 211,597 "username": "albanD",598 "name": "Alban D",599 "avatar_template": "/user_avatar/discuss.pytorch.org/alband/{size}/215_2.png",600 "post_count": 1,601 "primary_group_name": null,602 "flair_name": null,603 "flair_url": null,604 "flair_color": null,605 "flair_bg_color": null,606 "flair_group_id": null,607 "admin": true,608 "moderator": true,609 "trust_level": 4610 },611 {612 "id": 3534,613 "username": "ptrblck",614 "name": "",615 "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png",616 "post_count": 1,617 "primary_group_name": null,618 "flair_name": null,619 "flair_url": null,620 "flair_color": null,621 "flair_bg_color": null,622 "flair_group_id": null,623 "admin": true,624 "moderator": true,625 "trust_level": 2626 }627 ],628 "created_by": {629 "id": 24042,630 "username": "Shubhankar",631 "name": "",632 "avatar_template": "/letter_avatar_proxy/v4/letter/s/b9bd4f/{size}.png"633 },634 "last_poster": {635 "id": 3534,636 "username": "ptrblck",637 "name": "",638 "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png"639 },640 "links": [641 {642 "url": "https://github.com/tensorflow/tensorflow/issues/12871#issuecomment-370059148",643 "title": "About Deterministic Behaviour of GPU implementation of tensorflow · Issue #12871 · tensorflow/tensorflow · GitHub",644 "internal": false,645 "attachment": false,646 "reflection": false,647 "clicks": 10,648 "user_id": 24042,649 "domain": "github.com",650 "root_domain": "github.com"651 }652 ]653 },654 "bookmarks": []655 },656 {657 "post_stream": {658 "posts": [659 {660 "id": 154234,661 "name": "Fupanbo",662 "username": "fupanbo",663 "avatar_template": "/letter_avatar_proxy/v4/letter/f/ecae2f/{size}.png",664 "created_at": "2019-12-21T07:19:23.524Z",665 "cooked": "<p>so sorry for my pool english and thank u for watching my question</p>\n<p>i set batchsize=2<br>\nso after nn.Linear i get a 2x24 vector to use for my Classification task1,<br>\nnow i want to get a 2X8 vector from the 2X24 vector for my Classification task2,<br>\ni used split cat sum and so on to achieve the purpose<br>\nbut when i run the code ,It can run for a while and then report an error<br>\n{{{<br>\nTraceback (most recent call last):<br>\nFile “FeedbackNet_train.py”, line 65, in <br>\ntrainer.train()<br>\nFile “/home/fp/feedback/pytorch_feedback-network-master/utils/Trainer.py”, line 113, in train<br>\nself._train_one_epoch()<br>\nFile “/home/fp/feedback/pytorch_feedback-network-master/utils/Trainer.py”, line 184, in _train_one_epoch<br>\noutputs1,outputs2 = self.model(inputs) #<span class=\"hashtag\">#4x2x24</span><br>\nFile “/home/fp/.conda/envs/toold/lib/python2.7/site-packages/torch/nn/modules/module.py”, line 357, in <strong>call</strong><br>\nresult = self.forward(*input, **kwargs)<br>\nFile “/home/fp/.conda/envs/toold/lib/python2.7/site-packages/torch/nn/parallel/data_parallel.py”, line 71, in forward<br>\nreturn self.module(*inputs[0], **kwargs[0])<br>\nFile “/home/fp/.conda/envs/toold/lib/python2.7/site-packages/torch/nn/modules/module.py”, line 357, in <strong>call</strong><br>\nresult = self.forward(*input, **kwargs)<br>\nFile “/home/fp/feedback/pytorch_feedback-network-master/network/feedbacknet.py”, line 75, in forward<br>\nb1,b2=x_i.split(1,0)<br>\nValueError: need more than 1 value to unpack<br>\n}}}</p>\n<p>here is my code in network forward<br>\nx_finished = []<br>\nmx_finished = []</p>\n<pre><code> for x_i in x_all:\n x_i = F.relu(x_i) \n x_i =self.avg_pool(x_i) \n x_i=x_i.view(x.size()[0],-1)\n x_i=self.output(x_i) \n \n x_finished.append(x_i)\n \n x_i=self.soft(x_i) ''now x_i is a 2x24 vector ''\n\n b1,b2=x_i.split(1,0)\n b1=b1*w \"w is a constant tensor size 8X24 to help me transform the 1X24 to 1x8 \"\n b2=b2*w ’‘after b1=b1*w b1 from 1x24 to 8X24\"\"\n b1=torch.sum(b1,1) \"\"b1: 8x24 -> 8\"\"\n b1=b1.view(1,-1) ''b1 : 8 -> 1x8''\n b2=torch.sum(b2,1)\n b2=b2.view(1,-1)\n x_i=torch.cat((b1,b2),0) \"\"get x_i 2x8 size''\n mx_finished.append(x_i)\n\n return x_finished,mx_finished\n</code></pre>\n<p>because my network is based on CONVLSTM so x_finished have four 2x24<br>\nmx_finished have four 2x8<br>\nI think my code logic is smooth, but I don’t know why I get an error</p>",666 "post_number": 1,667 "post_type": 1,668 "posts_count": 10,669 "updated_at": "2019-12-21T07:26:09.766Z",670 "reply_count": 0,671 "reply_to_post_number": null,672 "quote_count": 0,673 "incoming_link_count": 164,674 "reads": 11,675 "readers_count": 10,676 "score": 822.2,677 "yours": false,678 "topic_id": 64697,679 "topic_slug": "valueerror-need-more-than-1-value-to-unpack",680 "display_username": "Fupanbo",681 "primary_group_name": null,682 "flair_name": null,683 "flair_url": null,684 "flair_bg_color": null,685 "flair_color": null,686 "flair_group_id": null,687 "badges_granted": [],688 "version": 2,689 "can_edit": false,690 "can_delete": false,691 "can_recover": false,692 "can_see_hidden_post": false,693 "can_wiki": false,694 "read": true,695 "user_title": null,696 "bookmarked": false,697 "actions_summary": [],698 "moderator": false,699 "admin": false,700 "staff": false,701 "user_id": 25802,702 "hidden": false,703 "trust_level": 1,704 "deleted_at": null,705 "user_deleted": false,706 "edit_reason": null,707 "can_view_edit_history": true,708 "wiki": false,709 "post_url": "/t/valueerror-need-more-than-1-value-to-unpack/64697/1",710 "can_accept_answer": false,711 "can_unaccept_answer": false,712 "accepted_answer": false,713 "topic_accepted_answer": null,714 "can_vote": false715 },716 {717 "id": 154310,718 "name": "",719 "username": "ptrblck",720 "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png",721 "created_at": "2019-12-21T20:31:30.001Z",722 "cooked": "<p>I guess the last batch could be smaller than the other ones, thus <code>x_i</code> might be e.g. <code>[1, 24]</code> instead of the expected <code>[2, 24]</code>.<br>\nCould you add a print statement right before the <code>.split</code> call and check the shapes during training, and which shape <code>x_i</code> has before throwing this error?</p>\n<p>A workaround could be to drop the last smaller batch via <code>drop_last=True</code> in your <code>DataLoader</code>.</p>",723 "post_number": 2,724 "post_type": 1,725 "posts_count": 10,726 "updated_at": "2019-12-21T20:31:30.001Z",727 "reply_count": 1,728 "reply_to_post_number": null,729 "quote_count": 0,730 "incoming_link_count": 1,731 "reads": 9,732 "readers_count": 8,733 "score": 11.8,734 "yours": false,735 "topic_id": 64697,736 "topic_slug": "valueerror-need-more-than-1-value-to-unpack",737 "display_username": "",738 "primary_group_name": null,739 "flair_name": null,740 "flair_url": null,741 "flair_bg_color": null,742 "flair_color": null,743 "flair_group_id": null,744 "badges_granted": [],745 "version": 1,746 "can_edit": false,747 "can_delete": false,748 "can_recover": false,749 "can_see_hidden_post": false,750 "can_wiki": false,751 "read": true,752 "user_title": "",753 "bookmarked": false,754 "actions_summary": [],755 "moderator": true,756 "admin": true,757 "staff": true,758 "user_id": 3534,759 "hidden": false,760 "trust_level": 2,761 "deleted_at": null,762 "user_deleted": false,763 "edit_reason": null,764 "can_view_edit_history": true,765 "wiki": false,766 "post_url": "/t/valueerror-need-more-than-1-value-to-unpack/64697/2",767 "can_accept_answer": false,768 "can_unaccept_answer": false,769 "accepted_answer": false,770 "topic_accepted_answer": null771 },772 {773 "id": 154350,774 "name": "Fupanbo",775 "username": "fupanbo",776 "avatar_template": "/letter_avatar_proxy/v4/letter/f/ecae2f/{size}.png",777 "created_at": "2019-12-22T08:19:33.964Z",778 "cooked": "<p>so thanks to u<br>\nu are right , the error caused by the last batch could be smaller than the other ones<br>\non the other hand i found some new questions want to ask u,<br>\nin my code the b1,b2 is the ‘’ for loop \" s local variable,<br>\nshould i set them as the Variable tensor in cuda before my network is working?<br>\ne.g. : b1=torch.randn(2,24)<br>\nb1=Variable(b1)<br>\nb1=b1.cuda()</p>\n<p>another question is:<br>\ni use split and cat to achieve “2x24 -->2X8”<br>\ni mean the split and cat seems like not a basic mathematical operations<br>\nso when the back propagation is running , it looks like the grad update will be effected?</p>\n<p>last one is:<br>\ni want to use the 2x24 vector get my loss1 and use the 2X8 vector get my loss2,<br>\nso the code for the part of \"2X24 turn to 2X8 \" should put in my network 's def forward<br>\nor in my every train epoch<br>\n(i mean the network just return the 2X24 and get 2X8 ,loss1,loss2 in the training<br>\nor the network return the 2X24 ,2X8 and get loss1,loss2 in the training )</p>\n<p>Thank you very much for helping me</p>",779 "post_number": 3,780 "post_type": 1,781 "posts_count": 10,782 "updated_at": "2019-12-22T08:19:33.964Z",783 "reply_count": 1,784 "reply_to_post_number": 2,785 "quote_count": 0,786 "incoming_link_count": 0,787 "reads": 10,788 "readers_count": 9,789 "score": 7.0,790 "yours": false,791 "topic_id": 64697,792 "topic_slug": "valueerror-need-more-than-1-value-to-unpack",793 "display_username": "Fupanbo",794 "primary_group_name": null,795 "flair_name": null,796 "flair_url": null,797 "flair_bg_color": null,798 "flair_color": null,799 "flair_group_id": null,800 "badges_granted": [],801 "version": 1,802 "can_edit": false,803 "can_delete": false,804 "can_recover": false,805 "can_see_hidden_post": false,806 "can_wiki": false,807 "read": true,808 "user_title": null,809 "reply_to_user": {810 "id": 3534,811 "username": "ptrblck",812 "name": "",813 "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png"814 },815 "bookmarked": false,816 "actions_summary": [],817 "moderator": false,818 "admin": false,819 "staff": false,820 "user_id": 25802,821 "hidden": false,822 "trust_level": 1,823 "deleted_at": null,824 "user_deleted": false,825 "edit_reason": null,826 "can_view_edit_history": true,827 "wiki": false,828 "post_url": "/t/valueerror-need-more-than-1-value-to-unpack/64697/3",829 "can_accept_answer": false,830 "can_unaccept_answer": false,831 "accepted_answer": false,832 "topic_accepted_answer": null833 },834 {835 "id": 154351,836 "name": "",837 "username": "ptrblck",838 "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png",839 "created_at": "2019-12-22T08:25:34.657Z",840 "cooked": "<aside class=\"quote no-group\" data-username=\"fupanbo\" data-post=\"3\" data-topic=\"64697\">\n<div class=\"title\">\n<div class=\"quote-controls\"></div>\n<img loading=\"lazy\" alt=\"\" width=\"24\" height=\"24\" src=\"https://discuss.pytorch.org/letter_avatar_proxy/v4/letter/f/ecae2f/48.png\" class=\"avatar\"> fupanbo:</div>\n<blockquote>\n<p>should i set them as the Variable</p>\n</blockquote>\n</aside>\n<p>No, <code>Variables</code> are deprecated since PyTorch <code>0.4.0</code>, so you can just use tensors in newer versions.</p>\n<aside class=\"quote no-group\" data-username=\"fupanbo\" data-post=\"3\" data-topic=\"64697\">\n<div class=\"title\">\n<div class=\"quote-controls\"></div>\n<img loading=\"lazy\" alt=\"\" width=\"24\" height=\"24\" src=\"https://discuss.pytorch.org/letter_avatar_proxy/v4/letter/f/ecae2f/48.png\" class=\"avatar\"> fupanbo:</div>\n<blockquote>\n<p>when the back propagation is running , it looks like the grad update will be effected?</p>\n</blockquote>\n</aside>\n<p><code>torch.split</code> and <code>torch.cat</code> won’t detach the tensors from the computation graph and the backward call will still work.</p>\n<aside class=\"quote no-group\" data-username=\"fupanbo\" data-post=\"3\" data-topic=\"64697\">\n<div class=\"title\">\n<div class=\"quote-controls\"></div>\n<img loading=\"lazy\" alt=\"\" width=\"24\" height=\"24\" src=\"https://discuss.pytorch.org/letter_avatar_proxy/v4/letter/f/ecae2f/48.png\" class=\"avatar\"> fupanbo:</div>\n<blockquote>\n<p>so the code for the part of \"2X24 turn to 2X8 \" should put in my network 's def forward<br>\nor in my every train epoch</p>\n</blockquote>\n</aside>\n<p>It doesn’t matter where these operations are called. As long as these ops are called somewhere, the results will be the same.<br>\nI would personally put these operations into the model and just return the output, which can be passed to a criterion, but your coding style might differ. <img src=\"https://discuss.pytorch.org/images/emoji/apple/wink.png?v=12\" title=\":wink:\" class=\"emoji\" alt=\":wink:\" loading=\"lazy\" width=\"20\" height=\"20\"></p>",841 "post_number": 4,842 "post_type": 1,843 "posts_count": 10,844 "updated_at": "2019-12-22T08:25:34.657Z",845 "reply_count": 1,846 "reply_to_post_number": 3,847 "quote_count": 1,848 "incoming_link_count": 4,849 "reads": 9,850 "readers_count": 8,851 "score": 26.8,852 "yours": false,853 "topic_id": 64697,854 "topic_slug": "valueerror-need-more-than-1-value-to-unpack",855 "display_username": "",856 "primary_group_name": null,857 "flair_name": null,858 "flair_url": null,859 "flair_bg_color": null,860 "flair_color": null,861 "flair_group_id": null,862 "badges_granted": [],863 "version": 1,864 "can_edit": false,865 "can_delete": false,866 "can_recover": false,867 "can_see_hidden_post": false,868 "can_wiki": false,869 "read": true,870 "user_title": "",871 "bookmarked": false,872 "actions_summary": [],873 "moderator": true,874 "admin": true,875 "staff": true,876 "user_id": 3534,877 "hidden": false,878 "trust_level": 2,879 "deleted_at": null,880 "user_deleted": false,881 "edit_reason": null,882 "can_view_edit_history": true,883 "wiki": false,884 "post_url": "/t/valueerror-need-more-than-1-value-to-unpack/64697/4",885 "can_accept_answer": false,886 "can_unaccept_answer": false,887 "accepted_answer": false,888 "topic_accepted_answer": null889 },890 {891 "id": 154352,892 "name": "Fupanbo",893 "username": "fupanbo",894 "avatar_template": "/letter_avatar_proxy/v4/letter/f/ecae2f/{size}.png",895 "created_at": "2019-12-22T08:48:54.213Z",896 "cooked": "<aside class=\"quote no-group\" data-username=\"ptrblck\" data-post=\"4\" data-topic=\"64697\">\n<div class=\"title\">\n<div class=\"quote-controls\"></div>\n<img loading=\"lazy\" alt=\"\" width=\"24\" height=\"24\" src=\"https://discuss.pytorch.org/user_avatar/discuss.pytorch.org/ptrblck/48/1823_2.png\" class=\"avatar\"> ptrblck:</div>\n<blockquote>\n<p>No, <code>Variables</code> are deprecated since PyTorch <code>0.4.0</code> , so you can just use tensors in newer versions.</p>\n</blockquote>\n</aside>\n<p>yeah i know and my study code is based on 0.3.0 now i have not much time to change the version:pleading_face:<br>\ni want to know outside \"the for loop \" should the b1,b2 need be defined as tensor such as:b1=torch.randn(1,24)</p>\n<p>or just as a local value in “the for loop”</p>\n<p>other question i already know what u mean so appreciated:kissing_heart:</p>",897 "post_number": 5,898 "post_type": 1,899 "posts_count": 10,900 "updated_at": "2019-12-22T08:48:54.213Z",901 "reply_count": 1,902 "reply_to_post_number": 4,903 "quote_count": 1,904 "incoming_link_count": 0,905 "reads": 7,906 "readers_count": 6,907 "score": 6.4,908 "yours": false,909 "topic_id": 64697,910 "topic_slug": "valueerror-need-more-than-1-value-to-unpack",911 "display_username": "Fupanbo",912 "primary_group_name": null,913 "flair_name": null,914 "flair_url": null,915 "flair_bg_color": null,916 "flair_color": null,917 "flair_group_id": null,918 "badges_granted": [],919 "version": 1,920 "can_edit": false,921 "can_delete": false,922 "can_recover": false,923 "can_see_hidden_post": false,924 "can_wiki": false,925 "read": true,926 "user_title": null,927 "bookmarked": false,928 "actions_summary": [],929 "moderator": false,930 "admin": false,931 "staff": false,932 "user_id": 25802,933 "hidden": false,934 "trust_level": 1,935 "deleted_at": null,936 "user_deleted": false,937 "edit_reason": null,938 "can_view_edit_history": true,939 "wiki": false,940 "post_url": "/t/valueerror-need-more-than-1-value-to-unpack/64697/5",941 "can_accept_answer": false,942 "can_unaccept_answer": false,943 "accepted_answer": false,944 "topic_accepted_answer": null945 },946 {947 "id": 154422,948 "name": "",949 "username": "ptrblck",950 "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png",951 "created_at": "2019-12-22T23:33:37.921Z",952 "cooked": "<p>You don’t need to define <code>b1</code> and <code>b2</code> before the for loop, if you create them via the <code>split</code> inside the loop.</p>",953 "post_number": 6,954 "post_type": 1,955 "posts_count": 10,956 "updated_at": "2019-12-23T02:59:34.669Z",957 "reply_count": 1,958 "reply_to_post_number": 5,959 "quote_count": 0,960 "incoming_link_count": 0,961 "reads": 6,962 "readers_count": 5,963 "score": 6.2,964 "yours": false,965 "topic_id": 64697,966 "topic_slug": "valueerror-need-more-than-1-value-to-unpack",967 "display_username": "",968 "primary_group_name": null,969 "flair_name": null,970 "flair_url": null,971 "flair_bg_color": null,972 "flair_color": null,973 "flair_group_id": null,974 "badges_granted": [],975 "version": 1,976 "can_edit": false,977 "can_delete": false,978 "can_recover": false,979 "can_see_hidden_post": false,980 "can_wiki": false,981 "read": true,982 "user_title": "",983 "reply_to_user": {984 "id": 25802,985 "username": "fupanbo",986 "name": "Fupanbo",987 "avatar_template": "/letter_avatar_proxy/v4/letter/f/ecae2f/{size}.png"988 },989 "bookmarked": false,990 "actions_summary": [],991 "moderator": true,992 "admin": true,993 "staff": true,994 "user_id": 3534,995 "hidden": false,996 "trust_level": 2,997 "deleted_at": null,998 "user_deleted": false,999 "edit_reason": null,1000 "can_view_edit_history": true,1001 "wiki": false,1002 "post_url": "/t/valueerror-need-more-than-1-value-to-unpack/64697/6",1003 "can_accept_answer": false,1004 "can_unaccept_answer": false,1005 "accepted_answer": false,1006 "topic_accepted_answer": null1007 },1008 {1009 "id": 154437,1010 "name": "Fupanbo",1011 "username": "fupanbo",1012 "avatar_template": "/letter_avatar_proxy/v4/letter/f/ecae2f/{size}.png",1013 "created_at": "2019-12-23T02:59:55.293Z",1014 "cooked": "<p>thank u<br>\ni recently read some literature their code is based on pytorch 0.4.0 or 0.4.1<br>\ndo u have some suggestions about the chioce between pytorch 0.4 and 1.0 or newer versions?</p>",1015 "post_number": 8,1016 "post_type": 1,1017 "posts_count": 10,1018 "updated_at": "2019-12-23T02:59:55.293Z",1019 "reply_count": 1,1020 "reply_to_post_number": 6,1021 "quote_count": 0,1022 "incoming_link_count": 1,1023 "reads": 5,1024 "readers_count": 4,1025 "score": 11.0,1026 "yours": false,1027 "topic_id": 64697,1028 "topic_slug": "valueerror-need-more-than-1-value-to-unpack",1029 "display_username": "Fupanbo",1030 "primary_group_name": null,1031 "flair_name": null,1032 "flair_url": null,1033 "flair_bg_color": null,1034 "flair_color": null,1035 "flair_group_id": null,1036 "badges_granted": [],1037 "version": 1,1038 "can_edit": false,1039 "can_delete": false,1040 "can_recover": false,1041 "can_see_hidden_post": false,1042 "can_wiki": false,1043 "read": true,1044 "user_title": null,1045 "reply_to_user": {1046 "id": 3534,1047 "username": "ptrblck",1048 "name": "",1049 "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png"1050 },1051 "bookmarked": false,1052 "actions_summary": [],1053 "moderator": false,1054 "admin": false,1055 "staff": false,1056 "user_id": 25802,1057 "hidden": false,1058 "trust_level": 1,1059 "deleted_at": null,1060 "user_deleted": false,1061 "edit_reason": null,1062 "can_view_edit_history": true,1063 "wiki": false,1064 "post_url": "/t/valueerror-need-more-than-1-value-to-unpack/64697/8",1065 "can_accept_answer": false,1066 "can_unaccept_answer": false,1067 "accepted_answer": false,1068 "topic_accepted_answer": null1069 },1070 {1071 "id": 154448,1072 "name": "",1073 "username": "ptrblck",1074 "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png",1075 "created_at": "2019-12-23T05:56:33.766Z",1076 "cooked": "<p>I would strongly recommend to use the latest stable release, which is <code>1.3.1</code> at the moment.<br>\nOtherwise you might run into legacy issues, which were already solved. Also, you’ll get all new features.</p>",1077 "post_number": 9,1078 "post_type": 1,1079 "posts_count": 10,1080 "updated_at": "2019-12-23T05:56:33.766Z",1081 "reply_count": 1,1082 "reply_to_post_number": 8,1083 "quote_count": 0,1084 "incoming_link_count": 1,1085 "reads": 5,1086 "readers_count": 4,1087 "score": 11.0,1088 "yours": false,1089 "topic_id": 64697,1090 "topic_slug": "valueerror-need-more-than-1-value-to-unpack",1091 "display_username": "",1092 "primary_group_name": null,1093 "flair_name": null,1094 "flair_url": null,1095 "flair_bg_color": null,1096 "flair_color": null,1097 "flair_group_id": null,1098 "badges_granted": [],1099 "version": 1,1100 "can_edit": false,1101 "can_delete": false,1102 "can_recover": false,1103 "can_see_hidden_post": false,1104 "can_wiki": false,1105 "read": true,1106 "user_title": "",1107 "reply_to_user": {1108 "id": 25802,1109 "username": "fupanbo",1110 "name": "Fupanbo",1111 "avatar_template": "/letter_avatar_proxy/v4/letter/f/ecae2f/{size}.png"1112 },1113 "bookmarked": false,1114 "actions_summary": [],1115 "moderator": true,1116 "admin": true,1117 "staff": true,1118 "user_id": 3534,1119 "hidden": false,1120 "trust_level": 2,1121 "deleted_at": null,1122 "user_deleted": false,1123 "edit_reason": null,1124 "can_view_edit_history": true,1125 "wiki": false,1126 "post_url": "/t/valueerror-need-more-than-1-value-to-unpack/64697/9",1127 "can_accept_answer": false,1128 "can_unaccept_answer": false,1129 "accepted_answer": false,1130 "topic_accepted_answer": null1131 },1132 {1133 "id": 154465,1134 "name": "Fupanbo",1135 "username": "fupanbo",1136 "avatar_template": "/letter_avatar_proxy/v4/letter/f/ecae2f/{size}.png",1137 "created_at": "2019-12-23T06:35:57.269Z",1138 "cooked": "<p>ok as a college student now i get ready for my new trip to e study pytorch<br>\nthank u for your help and patience all the time~</p>",1139 "post_number": 10,1140 "post_type": 1,1141 "posts_count": 10,1142 "updated_at": "2019-12-23T06:35:57.269Z",1143 "reply_count": 1,1144 "reply_to_post_number": 9,1145 "quote_count": 0,1146 "incoming_link_count": 7,1147 "reads": 5,1148 "readers_count": 4,1149 "score": 41.0,1150 "yours": false,1151 "topic_id": 64697,1152 "topic_slug": "valueerror-need-more-than-1-value-to-unpack",1153 "display_username": "Fupanbo",1154 "primary_group_name": null,1155 "flair_name": null,1156 "flair_url": null,1157 "flair_bg_color": null,1158 "flair_color": null,1159 "flair_group_id": null,1160 "badges_granted": [],1161 "version": 1,1162 "can_edit": false,1163 "can_delete": false,1164 "can_recover": false,1165 "can_see_hidden_post": false,1166 "can_wiki": false,1167 "read": true,1168 "user_title": null,1169 "reply_to_user": {1170 "id": 3534,1171 "username": "ptrblck",1172 "name": "",1173 "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png"1174 },1175 "bookmarked": false,1176 "actions_summary": [],1177 "moderator": false,1178 "admin": false,1179 "staff": false,1180 "user_id": 25802,1181 "hidden": false,1182 "trust_level": 1,1183 "deleted_at": null,1184 "user_deleted": false,1185 "edit_reason": null,1186 "can_view_edit_history": true,1187 "wiki": false,1188 "post_url": "/t/valueerror-need-more-than-1-value-to-unpack/64697/10",1189 "can_accept_answer": false,1190 "can_unaccept_answer": false,1191 "accepted_answer": false,1192 "topic_accepted_answer": null1193 },1194 {1195 "id": 154468,1196 "name": "",1197 "username": "ptrblck",1198 "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png",1199 "created_at": "2019-12-23T06:40:14.240Z",1200 "cooked": "<p>Sure! Feel free to post your questions in this board (or search for related questions), in case you get stuck. <img src=\"https://discuss.pytorch.org/images/emoji/apple/wink.png?v=9\" title=\":wink:\" class=\"emoji\" alt=\":wink:\"></p>",