CoolFace
Datasetpublic

Anurag1734/cuda-error-resolution-analysis

sourceHugging Faceupdated 2mo agoView on Hugging Face
0likes7downloads
topics_batch_538.json61955 linesDownload Raw Back to raw
1[2  {3    "post_stream": {4      "posts": [5        {6          "id": 154159,7          "name": "",8          "username": "Shubhankar",9          "avatar_template": "/letter_avatar_proxy/v4/letter/s/b9bd4f/{size}.png",10          "created_at": "2019-12-20T20:39:56.176Z",11          "cooked": "<p>How to use NVIDIA AUTO TUNE with pytorch.</p>",12          "post_number": 1,13          "post_type": 1,14          "posts_count": 4,15          "updated_at": "2019-12-20T20:39:56.176Z",16          "reply_count": 0,17          "reply_to_post_number": null,18          "quote_count": 0,19          "incoming_link_count": 745,20          "reads": 26,21          "readers_count": 25,22          "score": 3725.2,23          "yours": false,24          "topic_id": 64669,25          "topic_slug": "how-to-use-auto-tune-with-pytorch",26          "display_username": "",27          "primary_group_name": null,28          "flair_name": null,29          "flair_url": null,30          "flair_bg_color": null,31          "flair_color": null,32          "flair_group_id": null,33          "badges_granted": [],34          "version": 1,35          "can_edit": false,36          "can_delete": false,37          "can_recover": false,38          "can_see_hidden_post": false,39          "can_wiki": false,40          "read": true,41          "user_title": "",42          "bookmarked": false,43          "actions_summary": [],44          "moderator": false,45          "admin": false,46          "staff": false,47          "user_id": 24042,48          "hidden": false,49          "trust_level": 1,50          "deleted_at": null,51          "user_deleted": false,52          "edit_reason": null,53          "can_view_edit_history": true,54          "wiki": false,55          "post_url": "/t/how-to-use-auto-tune-with-pytorch/64669/1",56          "can_accept_answer": false,57          "can_unaccept_answer": false,58          "accepted_answer": false,59          "topic_accepted_answer": true,60          "can_vote": false61        },62        {63          "id": 154388,64          "name": "Alban D",65          "username": "albanD",66          "avatar_template": "/user_avatar/discuss.pytorch.org/alband/{size}/215_2.png",67          "created_at": "2019-12-22T13:47:57.180Z",68          "cooked": "<p>Hi,</p>\n<p>I am not familiar with this? Could you explain what it is and what it does please?</p>",69          "post_number": 2,70          "post_type": 1,71          "posts_count": 4,72          "updated_at": "2019-12-22T13:47:57.180Z",73          "reply_count": 0,74          "reply_to_post_number": null,75          "quote_count": 0,76          "incoming_link_count": 5,77          "reads": 24,78          "readers_count": 23,79          "score": 29.8,80          "yours": false,81          "topic_id": 64669,82          "topic_slug": "how-to-use-auto-tune-with-pytorch",83          "display_username": "Alban D",84          "primary_group_name": null,85          "flair_name": null,86          "flair_url": null,87          "flair_bg_color": null,88          "flair_color": null,89          "flair_group_id": null,90          "badges_granted": [],91          "version": 1,92          "can_edit": false,93          "can_delete": false,94          "can_recover": false,95          "can_see_hidden_post": false,96          "can_wiki": false,97          "read": true,98          "user_title": "",99          "bookmarked": false,100          "actions_summary": [],101          "moderator": true,102          "admin": true,103          "staff": true,104          "user_id": 211,105          "hidden": false,106          "trust_level": 4,107          "deleted_at": null,108          "user_deleted": false,109          "edit_reason": null,110          "can_view_edit_history": true,111          "wiki": false,112          "post_url": "/t/how-to-use-auto-tune-with-pytorch/64669/2",113          "can_accept_answer": false,114          "can_unaccept_answer": false,115          "accepted_answer": false,116          "topic_accepted_answer": true117        },118        {119          "id": 154407,120          "name": "",121          "username": "Shubhankar",122          "avatar_template": "/letter_avatar_proxy/v4/letter/s/b9bd4f/{size}.png",123          "created_at": "2019-12-22T17:23:02.009Z",124          "cooked": "<aside class=\"onebox githubissue\">\n  <header class=\"source\">\n      <a href=\"https://github.com/tensorflow/tensorflow/issues/12871#issuecomment-370059148\" target=\"_blank\" rel=\"nofollow noopener\">github.com/tensorflow/tensorflow</a>\n  </header>\n  <article class=\"onebox-body\">\n    <div class=\"github-row\">\n  <div class=\"github-icon-container\" title=\"Issue\">\n\t  <svg width=\"60\" height=\"60\" class=\"github-icon\" viewbox=\"0 0 14 16\" aria-hidden=\"true\"><path d=\"M7 2.3c3.14 0 5.7 2.56 5.7 5.7s-2.56 5.7-5.7 5.7A5.71 5.71 0 0 1 1.3 8c0-3.14 2.56-5.7 5.7-5.7zM7 1C3.14 1 0 4.14 0 8s3.14 7 7 7 7-3.14 7-7-3.14-7-7-7zm1 3H6v5h2V4zm0 6H6v2h2v-2z\"></path></svg>\n  </div>\n\n  <div class=\"github-info-container\">\n    <h4>\n      <a href=\"https://github.com/tensorflow/tensorflow/issues/12871#issuecomment-370059148\" target=\"_blank\" rel=\"nofollow noopener\">About Deterministic Behaviour of GPU implementation of tensorflow</a>\n    </h4>\n\n    <div class=\"github-info\">\n      <div class=\"date\">\n        opened <span class=\"discourse-local-date\" data-format=\"ll\" data-date=\"2017-09-07\" data-time=\"08:07:09\" data-timezone=\"UTC\">08:07AM - 07 Sep 17 UTC</span>\n      </div>\n\n        <div class=\"date\">\n          closed <span class=\"discourse-local-date\" data-format=\"ll\" data-date=\"2018-09-21\" data-time=\"18:57:17\" data-timezone=\"UTC\">06:57PM - 21 Sep 18 UTC</span>\n        </div>\n\n      <div class=\"user\">\n        <a href=\"https://github.com/antares1987\" target=\"_blank\" rel=\"nofollow noopener\">\n          <img alt=\"antares1987\" src=\"https://avatars0.githubusercontent.com/u/31724739?v=4\" class=\"onebox-avatar-inline\" width=\"20\" height=\"20\">\n          antares1987\n        </a>\n      </div>\n    </div>\n  </div>\n</div>\n\n<div class=\"github-row\">\n  <p class=\"github-content\">OS Platform and Distribution (e.g., Linux Ubuntu 16.04): Debian\nTensorFlow version (use command below):('v1.3.0-rc2-20-g0787eee', '1.3.0')\nCUDA/cuDNN version: 8.0\nGPU model and memory: GeForce GTX...</p>\n</div>\n\n<div class=\"labels\">\n    <span style=\"display:inline-block;margin-top:2px;background-color: #B8B8B8;padding: 2px;border-radius: 4px;color: #fff;margin-left: 3px;\">stat:awaiting response</span>\n</div>\n\n  </article>\n  <div class=\"onebox-metadata\">\n    \n    \n  </div>\n  <div style=\"clear: both\"></div>\n</aside>\n",125          "post_number": 3,126          "post_type": 1,127          "posts_count": 4,128          "updated_at": "2019-12-22T17:23:02.009Z",129          "reply_count": 1,130          "reply_to_post_number": null,131          "quote_count": 0,132          "incoming_link_count": 5,133          "reads": 23,134          "readers_count": 22,135          "score": 34.6,136          "yours": false,137          "topic_id": 64669,138          "topic_slug": "how-to-use-auto-tune-with-pytorch",139          "display_username": "",140          "primary_group_name": null,141          "flair_name": null,142          "flair_url": null,143          "flair_bg_color": null,144          "flair_color": null,145          "flair_group_id": null,146          "badges_granted": [],147          "version": 1,148          "can_edit": false,149          "can_delete": false,150          "can_recover": false,151          "can_see_hidden_post": false,152          "can_wiki": false,153          "link_counts": [154            {155              "url": "https://github.com/tensorflow/tensorflow/issues/12871#issuecomment-370059148",156              "internal": false,157              "reflection": false,158              "title": "About Deterministic Behaviour of GPU implementation of tensorflow · Issue #12871 · tensorflow/tensorflow · GitHub",159              "clicks": 10160            },161            {162              "url": "https://github.com/antares1987",163              "internal": false,164              "reflection": false,165              "title": "antares1987 · GitHub",166              "clicks": 0167            }168          ],169          "read": true,170          "user_title": "",171          "bookmarked": false,172          "actions_summary": [],173          "moderator": false,174          "admin": false,175          "staff": false,176          "user_id": 24042,177          "hidden": false,178          "trust_level": 1,179          "deleted_at": null,180          "user_deleted": false,181          "edit_reason": null,182          "can_view_edit_history": true,183          "wiki": false,184          "post_url": "/t/how-to-use-auto-tune-with-pytorch/64669/3",185          "can_accept_answer": false,186          "can_unaccept_answer": false,187          "accepted_answer": false,188          "topic_accepted_answer": true189        },190        {191          "id": 154475,192          "name": "",193          "username": "ptrblck",194          "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png",195          "created_at": "2019-12-23T07:04:18.869Z",196          "cooked": "<p>I’m not familiar with TensorFlow, but it seems their “auto tune” functionality corresponds to <code>torch.backends.cudnn.benchmark = True</code> mode.</p>",197          "post_number": 4,198          "post_type": 1,199          "posts_count": 4,200          "updated_at": "2020-01-06T05:18:09.159Z",201          "reply_count": 0,202          "reply_to_post_number": 3,203          "quote_count": 0,204          "incoming_link_count": 7,205          "reads": 19,206          "readers_count": 18,207          "score": 83.8,208          "yours": false,209          "topic_id": 64669,210          "topic_slug": "how-to-use-auto-tune-with-pytorch",211          "display_username": "",212          "primary_group_name": null,213          "flair_name": null,214          "flair_url": null,215          "flair_bg_color": null,216          "flair_color": null,217          "flair_group_id": null,218          "badges_granted": [],219          "version": 1,220          "can_edit": false,221          "can_delete": false,222          "can_recover": false,223          "can_see_hidden_post": false,224          "can_wiki": false,225          "read": true,226          "user_title": "",227          "reply_to_user": {228            "id": 24042,229            "username": "Shubhankar",230            "name": "",231            "avatar_template": "/letter_avatar_proxy/v4/letter/s/b9bd4f/{size}.png"232          },233          "bookmarked": false,234          "actions_summary": [235            {236              "id": 2,237              "count": 1238            }239          ],240          "moderator": true,241          "admin": true,242          "staff": true,243          "user_id": 3534,244          "hidden": false,245          "trust_level": 2,246          "deleted_at": null,247          "user_deleted": false,248          "edit_reason": null,249          "can_view_edit_history": true,250          "wiki": false,251          "post_url": "/t/how-to-use-auto-tune-with-pytorch/64669/4",252          "can_accept_answer": false,253          "can_unaccept_answer": false,254          "accepted_answer": true,255          "topic_accepted_answer": true256        }257      ],258      "stream": [259        154159,260        154388,261        154407,262        154475263      ]264    },265    "timeline_lookup": [266      [267        1,268        2136269      ],270      [271        2,272        2134273      ]274    ],275    "suggested_topics": [276      {277        "fancy_title": "ExecuTorch: Variable-Length Inputs for Export and F.maxpool1d",278        "id": 218652,279        "title": "ExecuTorch: Variable-Length Inputs for Export and F.maxpool1d",280        "slug": "executorch-variable-length-inputs-for-export-and-f-maxpool1d",281        "posts_count": 1,282        "reply_count": 0,283        "highest_post_number": 1,284        "image_url": null,285        "created_at": "2025-04-06T02:00:56.155Z",286        "last_posted_at": "2025-04-06T02:00:56.195Z",287        "bumped": true,288        "bumped_at": "2025-04-06T02:00:56.195Z",289        "archetype": "regular",290        "unseen": false,291        "pinned": false,292        "unpinned": null,293        "visible": true,294        "closed": false,295        "archived": false,296        "bookmarked": null,297        "liked": null,298        "tags_descriptions": {},299        "like_count": 0,300        "views": 19,301        "category_id": 1,302        "featured_link": null,303        "has_accepted_answer": false,304        "posters": [305          {306            "extras": "latest single",307            "description": "Original Poster, Most Recent Poster",308            "user": {309              "id": 83653,310              "username": "cba913",311              "name": "",312              "avatar_template": "/user_avatar/discuss.pytorch.org/cba913/{size}/76501_2.png",313              "trust_level": 0314            }315          }316        ]317      },318      {319        "fancy_title": "How to approach a real-life problem while using rainfall data",320        "id": 212728,321        "title": "How to approach a real-life problem while using rainfall data",322        "slug": "how-to-approach-a-real-life-problem-while-using-rainfall-data",323        "posts_count": 1,324        "reply_count": 0,325        "highest_post_number": 1,326        "image_url": null,327        "created_at": "2024-11-09T07:21:57.841Z",328        "last_posted_at": "2024-11-09T07:21:57.888Z",329        "bumped": true,330        "bumped_at": "2024-11-09T07:21:57.888Z",331        "archetype": "regular",332        "unseen": false,333        "pinned": false,334        "unpinned": null,335        "visible": true,336        "closed": false,337        "archived": false,338        "bookmarked": null,339        "liked": null,340        "tags_descriptions": {},341        "like_count": 0,342        "views": 25,343        "category_id": 1,344        "featured_link": null,345        "has_accepted_answer": false,346        "posters": [347          {348            "extras": "latest single",349            "description": "Original Poster, Most Recent Poster",350            "user": {351              "id": 80781,352              "username": "Ritam_Pradhan",353              "name": "Ritam Pradhan",354              "avatar_template": "/user_avatar/discuss.pytorch.org/ritam_pradhan/{size}/73876_2.png",355              "trust_level": 1356            }357          }358        ]359      },360      {361        "fancy_title": "CUDA out of memory while using Llama3.1-8B for inference",362        "id": 213808,363        "title": "CUDA out of memory while using Llama3.1-8B for inference",364        "slug": "cuda-out-of-memory-while-using-llama3-1-8b-for-inference",365        "posts_count": 1,366        "reply_count": 0,367        "highest_post_number": 1,368        "image_url": null,369        "created_at": "2024-12-04T17:52:39.019Z",370        "last_posted_at": "2024-12-04T17:52:39.070Z",371        "bumped": true,372        "bumped_at": "2024-12-04T17:52:39.070Z",373        "archetype": "regular",374        "unseen": false,375        "pinned": false,376        "unpinned": null,377        "visible": true,378        "closed": false,379        "archived": false,380        "bookmarked": null,381        "liked": null,382        "tags_descriptions": {},383        "like_count": 0,384        "views": 96,385        "category_id": 1,386        "featured_link": null,387        "has_accepted_answer": false,388        "posters": [389          {390            "extras": "latest single",391            "description": "Original Poster, Most Recent Poster",392            "user": {393              "id": 70667,394              "username": "tomwagstaff-opml",395              "name": "Tom Wagstaff",396              "avatar_template": "/user_avatar/discuss.pytorch.org/tomwagstaff-opml/{size}/65211_2.png",397              "trust_level": 0398            }399          }400        ]401      },402      {403        "fancy_title": "ROCm: hipBLASLt error with gfx1103",404        "id": 212721,405        "title": "ROCm: hipBLASLt error with gfx1103",406        "slug": "rocm-hipblaslt-error-with-gfx1103",407        "posts_count": 2,408        "reply_count": 0,409        "highest_post_number": 2,410        "image_url": null,411        "created_at": "2024-11-09T01:43:15.446Z",412        "last_posted_at": "2025-03-03T19:42:45.282Z",413        "bumped": true,414        "bumped_at": "2025-03-03T19:42:45.282Z",415        "archetype": "regular",416        "unseen": false,417        "pinned": false,418        "unpinned": null,419        "visible": true,420        "closed": false,421        "archived": false,422        "bookmarked": null,423        "liked": null,424        "tags_descriptions": {},425        "like_count": 0,426        "views": 1070,427        "category_id": 1,428        "featured_link": null,429        "has_accepted_answer": false,430        "posters": [431          {432            "extras": null,433            "description": "Original Poster",434            "user": {435              "id": 78358,436              "username": "lumie",437              "name": "",438              "avatar_template": "/letter_avatar_proxy/v4/letter/l/ea5d25/{size}.png",439              "trust_level": 1440            }441          },442          {443            "extras": "latest",444            "description": "Most Recent Poster",445            "user": {446              "id": 82743,447              "username": "fngarrett",448              "name": "Garrett",449              "avatar_template": "/letter_avatar_proxy/v4/letter/f/e79b87/{size}.png",450              "trust_level": 1451            }452          }453        ]454      },455      {456        "fancy_title": "Tuning a network, subset of data, one non-frozen layer",457        "id": 218788,458        "title": "Tuning a network, subset of data, one non-frozen layer",459        "slug": "tuning-a-network-subset-of-data-one-non-frozen-layer",460        "posts_count": 1,461        "reply_count": 0,462        "highest_post_number": 1,463        "image_url": null,464        "created_at": "2025-04-07T05:54:35.100Z",465        "last_posted_at": "2025-04-07T05:54:35.140Z",466        "bumped": true,467        "bumped_at": "2025-04-07T05:54:35.140Z",468        "archetype": "regular",469        "unseen": false,470        "pinned": false,471        "unpinned": null,472        "visible": true,473        "closed": false,474        "archived": false,475        "bookmarked": null,476        "liked": null,477        "tags_descriptions": {},478        "like_count": 0,479        "views": 18,480        "category_id": 1,481        "featured_link": null,482        "has_accepted_answer": false,483        "posters": [484          {485            "extras": "latest single",486            "description": "Original Poster, Most Recent Poster",487            "user": {488              "id": 78333,489              "username": "emerth",490              "name": "",491              "avatar_template": "/user_avatar/discuss.pytorch.org/emerth/{size}/75379_2.png",492              "trust_level": 1493            }494          }495        ]496      }497    ],498    "tags_descriptions": {},499    "fancy_title": "How to use AUTO_TUNE with pytorch?",500    "id": 64669,501    "title": "How to use AUTO_TUNE with pytorch?",502    "posts_count": 4,503    "created_at": "2019-12-20T20:39:56.126Z",504    "views": 1525,505    "reply_count": 1,506    "like_count": 1,507    "last_posted_at": "2019-12-23T07:04:18.869Z",508    "visible": true,509    "closed": false,510    "archived": false,511    "has_summary": false,512    "archetype": "regular",513    "slug": "how-to-use-auto-tune-with-pytorch",514    "category_id": 1,515    "word_count": 56,516    "deleted_at": null,517    "user_id": 24042,518    "featured_link": null,519    "pinned_globally": false,520    "pinned_at": null,521    "pinned_until": null,522    "image_url": null,523    "slow_mode_seconds": 0,524    "draft": null,525    "draft_key": "topic_64669",526    "draft_sequence": null,527    "unpinned": null,528    "pinned": false,529    "current_post_number": 1,530    "highest_post_number": 4,531    "deleted_by": null,532    "actions_summary": [533      {534        "id": 4,535        "count": 0,536        "hidden": false,537        "can_act": false538      },539      {540        "id": 8,541        "count": 0,542        "hidden": false,543        "can_act": false544      },545      {546        "id": 10,547        "count": 0,548        "hidden": false,549        "can_act": false550      },551      {552        "id": 7,553        "count": 0,554        "hidden": false,555        "can_act": false556      }557    ],558    "chunk_size": 20,559    "bookmarked": false,560    "topic_timer": null,561    "message_bus_last_id": 0,562    "participant_count": 3,563    "show_read_indicator": false,564    "thumbnails": null,565    "slow_mode_enabled_until": null,566    "accepted_answer": {567      "post_number": 4,568      "username": "ptrblck",569      "name": "",570      "excerpt": "I’m not familiar with TensorFlow, but it seems their “auto tune” functionality corresponds to torch.backends.cudnn.benchmark = True mode."571    },572    "can_vote": false,573    "vote_count": 0,574    "user_voted": false,575    "discourse_zendesk_plugin_zendesk_id": null,576    "discourse_zendesk_plugin_zendesk_url": "https://your-url.zendesk.com/agent/tickets/",577    "details": {578      "can_edit": false,579      "notification_level": 1,580      "participants": [581        {582          "id": 24042,583          "username": "Shubhankar",584          "name": "",585          "avatar_template": "/letter_avatar_proxy/v4/letter/s/b9bd4f/{size}.png",586          "post_count": 2,587          "primary_group_name": null,588          "flair_name": null,589          "flair_url": null,590          "flair_color": null,591          "flair_bg_color": null,592          "flair_group_id": null,593          "trust_level": 1594        },595        {596          "id": 211,597          "username": "albanD",598          "name": "Alban D",599          "avatar_template": "/user_avatar/discuss.pytorch.org/alband/{size}/215_2.png",600          "post_count": 1,601          "primary_group_name": null,602          "flair_name": null,603          "flair_url": null,604          "flair_color": null,605          "flair_bg_color": null,606          "flair_group_id": null,607          "admin": true,608          "moderator": true,609          "trust_level": 4610        },611        {612          "id": 3534,613          "username": "ptrblck",614          "name": "",615          "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png",616          "post_count": 1,617          "primary_group_name": null,618          "flair_name": null,619          "flair_url": null,620          "flair_color": null,621          "flair_bg_color": null,622          "flair_group_id": null,623          "admin": true,624          "moderator": true,625          "trust_level": 2626        }627      ],628      "created_by": {629        "id": 24042,630        "username": "Shubhankar",631        "name": "",632        "avatar_template": "/letter_avatar_proxy/v4/letter/s/b9bd4f/{size}.png"633      },634      "last_poster": {635        "id": 3534,636        "username": "ptrblck",637        "name": "",638        "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png"639      },640      "links": [641        {642          "url": "https://github.com/tensorflow/tensorflow/issues/12871#issuecomment-370059148",643          "title": "About Deterministic Behaviour of GPU implementation of tensorflow · Issue #12871 · tensorflow/tensorflow · GitHub",644          "internal": false,645          "attachment": false,646          "reflection": false,647          "clicks": 10,648          "user_id": 24042,649          "domain": "github.com",650          "root_domain": "github.com"651        }652      ]653    },654    "bookmarks": []655  },656  {657    "post_stream": {658      "posts": [659        {660          "id": 154234,661          "name": "Fupanbo",662          "username": "fupanbo",663          "avatar_template": "/letter_avatar_proxy/v4/letter/f/ecae2f/{size}.png",664          "created_at": "2019-12-21T07:19:23.524Z",665          "cooked": "<p>so sorry for my  pool english and thank u for watching my question</p>\n<p>i set batchsize=2<br>\nso after nn.Linear i  get a 2x24 vector to use for my Classification task1,<br>\nnow i want to get a 2X8 vector from the 2X24 vector  for my Classification task2,<br>\ni used split cat sum and so on to achieve the purpose<br>\nbut when i run the code ,It can run for a while and then report an error<br>\n{{{<br>\nTraceback (most recent call last):<br>\nFile “FeedbackNet_train.py”, line 65, in <br>\ntrainer.train()<br>\nFile “/home/fp/feedback/pytorch_feedback-network-master/utils/Trainer.py”, line 113, in train<br>\nself._train_one_epoch()<br>\nFile “/home/fp/feedback/pytorch_feedback-network-master/utils/Trainer.py”, line 184, in _train_one_epoch<br>\noutputs1,outputs2 = self.model(inputs)      #<span class=\"hashtag\">#4x2x24</span><br>\nFile “/home/fp/.conda/envs/toold/lib/python2.7/site-packages/torch/nn/modules/module.py”, line 357, in <strong>call</strong><br>\nresult = self.forward(*input, **kwargs)<br>\nFile “/home/fp/.conda/envs/toold/lib/python2.7/site-packages/torch/nn/parallel/data_parallel.py”, line 71, in forward<br>\nreturn self.module(*inputs[0], **kwargs[0])<br>\nFile “/home/fp/.conda/envs/toold/lib/python2.7/site-packages/torch/nn/modules/module.py”, line 357, in <strong>call</strong><br>\nresult = self.forward(*input, **kwargs)<br>\nFile “/home/fp/feedback/pytorch_feedback-network-master/network/feedbacknet.py”, line 75, in forward<br>\nb1,b2=x_i.split(1,0)<br>\nValueError: need more than 1 value to unpack<br>\n}}}</p>\n<p>here is my code in network forward<br>\nx_finished = []<br>\nmx_finished = []</p>\n<pre><code>    for x_i in x_all:\n        x_i = F.relu(x_i) \n        x_i =self.avg_pool(x_i)          \n        x_i=x_i.view(x.size()[0],-1)\n        x_i=self.output(x_i)           \n        \n        x_finished.append(x_i)\n         \n        x_i=self.soft(x_i)     ''now x_i is a 2x24 vector ''\n\n        b1,b2=x_i.split(1,0)\n        b1=b1*w               \"w is  a constant tensor size 8X24 to help me transform the 1X24 to 1x8 \"\n        b2=b2*w               ’‘after b1=b1*w   b1 from 1x24 to 8X24\"\"\n        b1=torch.sum(b1,1)                  \"\"b1: 8x24 -&gt; 8\"\"\n        b1=b1.view(1,-1)                      ''b1 : 8 -&gt; 1x8''\n        b2=torch.sum(b2,1)\n        b2=b2.view(1,-1)\n        x_i=torch.cat((b1,b2),0)                   \"\"get x_i  2x8 size''\n        mx_finished.append(x_i)\n\n    return x_finished,mx_finished\n</code></pre>\n<p>because my network is based on CONVLSTM  so x_finished  have four  2x24<br>\nmx_finished have four 2x8<br>\nI think my code logic is smooth, but I don’t know why I get an error</p>",666          "post_number": 1,667          "post_type": 1,668          "posts_count": 10,669          "updated_at": "2019-12-21T07:26:09.766Z",670          "reply_count": 0,671          "reply_to_post_number": null,672          "quote_count": 0,673          "incoming_link_count": 164,674          "reads": 11,675          "readers_count": 10,676          "score": 822.2,677          "yours": false,678          "topic_id": 64697,679          "topic_slug": "valueerror-need-more-than-1-value-to-unpack",680          "display_username": "Fupanbo",681          "primary_group_name": null,682          "flair_name": null,683          "flair_url": null,684          "flair_bg_color": null,685          "flair_color": null,686          "flair_group_id": null,687          "badges_granted": [],688          "version": 2,689          "can_edit": false,690          "can_delete": false,691          "can_recover": false,692          "can_see_hidden_post": false,693          "can_wiki": false,694          "read": true,695          "user_title": null,696          "bookmarked": false,697          "actions_summary": [],698          "moderator": false,699          "admin": false,700          "staff": false,701          "user_id": 25802,702          "hidden": false,703          "trust_level": 1,704          "deleted_at": null,705          "user_deleted": false,706          "edit_reason": null,707          "can_view_edit_history": true,708          "wiki": false,709          "post_url": "/t/valueerror-need-more-than-1-value-to-unpack/64697/1",710          "can_accept_answer": false,711          "can_unaccept_answer": false,712          "accepted_answer": false,713          "topic_accepted_answer": null,714          "can_vote": false715        },716        {717          "id": 154310,718          "name": "",719          "username": "ptrblck",720          "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png",721          "created_at": "2019-12-21T20:31:30.001Z",722          "cooked": "<p>I guess the last batch could be smaller than the other ones, thus <code>x_i</code> might be e.g. <code>[1, 24]</code> instead of the expected <code>[2, 24]</code>.<br>\nCould you add a print statement right before the <code>.split</code> call and check the shapes during training, and which shape <code>x_i</code> has before throwing this error?</p>\n<p>A workaround could be to drop the last smaller batch via <code>drop_last=True</code> in your <code>DataLoader</code>.</p>",723          "post_number": 2,724          "post_type": 1,725          "posts_count": 10,726          "updated_at": "2019-12-21T20:31:30.001Z",727          "reply_count": 1,728          "reply_to_post_number": null,729          "quote_count": 0,730          "incoming_link_count": 1,731          "reads": 9,732          "readers_count": 8,733          "score": 11.8,734          "yours": false,735          "topic_id": 64697,736          "topic_slug": "valueerror-need-more-than-1-value-to-unpack",737          "display_username": "",738          "primary_group_name": null,739          "flair_name": null,740          "flair_url": null,741          "flair_bg_color": null,742          "flair_color": null,743          "flair_group_id": null,744          "badges_granted": [],745          "version": 1,746          "can_edit": false,747          "can_delete": false,748          "can_recover": false,749          "can_see_hidden_post": false,750          "can_wiki": false,751          "read": true,752          "user_title": "",753          "bookmarked": false,754          "actions_summary": [],755          "moderator": true,756          "admin": true,757          "staff": true,758          "user_id": 3534,759          "hidden": false,760          "trust_level": 2,761          "deleted_at": null,762          "user_deleted": false,763          "edit_reason": null,764          "can_view_edit_history": true,765          "wiki": false,766          "post_url": "/t/valueerror-need-more-than-1-value-to-unpack/64697/2",767          "can_accept_answer": false,768          "can_unaccept_answer": false,769          "accepted_answer": false,770          "topic_accepted_answer": null771        },772        {773          "id": 154350,774          "name": "Fupanbo",775          "username": "fupanbo",776          "avatar_template": "/letter_avatar_proxy/v4/letter/f/ecae2f/{size}.png",777          "created_at": "2019-12-22T08:19:33.964Z",778          "cooked": "<p>so thanks to u<br>\nu are right , the error  caused by  the last batch could be smaller than the other ones<br>\non the other hand i found some new questions want to ask u,<br>\nin my code  the b1,b2 is the  ‘’ for loop \" s  local variable,<br>\nshould i set them  as the Variable tensor in cuda before my network is working?<br>\ne.g.  :  b1=torch.randn(2,24)<br>\nb1=Variable(b1)<br>\nb1=b1.cuda()</p>\n<p>another question  is:<br>\ni use split and cat  to achieve “2x24 --&gt;2X8”<br>\ni mean the split and cat seems like not a basic mathematical operations<br>\nso when the back propagation is running , it looks like the grad update will be effected?</p>\n<p>last one is:<br>\ni want to use the 2x24 vector get my loss1  and use the 2X8 vector get my loss2,<br>\nso the code for the part of  \"2X24 turn to 2X8 \" should put in my  network 's  def forward<br>\nor in my every train epoch<br>\n(i mean the network just return the 2X24 and get 2X8 ,loss1,loss2 in the training<br>\nor the network return the 2X24 ,2X8 and get loss1,loss2 in the training  )</p>\n<p>Thank you very much for helping me</p>",779          "post_number": 3,780          "post_type": 1,781          "posts_count": 10,782          "updated_at": "2019-12-22T08:19:33.964Z",783          "reply_count": 1,784          "reply_to_post_number": 2,785          "quote_count": 0,786          "incoming_link_count": 0,787          "reads": 10,788          "readers_count": 9,789          "score": 7.0,790          "yours": false,791          "topic_id": 64697,792          "topic_slug": "valueerror-need-more-than-1-value-to-unpack",793          "display_username": "Fupanbo",794          "primary_group_name": null,795          "flair_name": null,796          "flair_url": null,797          "flair_bg_color": null,798          "flair_color": null,799          "flair_group_id": null,800          "badges_granted": [],801          "version": 1,802          "can_edit": false,803          "can_delete": false,804          "can_recover": false,805          "can_see_hidden_post": false,806          "can_wiki": false,807          "read": true,808          "user_title": null,809          "reply_to_user": {810            "id": 3534,811            "username": "ptrblck",812            "name": "",813            "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png"814          },815          "bookmarked": false,816          "actions_summary": [],817          "moderator": false,818          "admin": false,819          "staff": false,820          "user_id": 25802,821          "hidden": false,822          "trust_level": 1,823          "deleted_at": null,824          "user_deleted": false,825          "edit_reason": null,826          "can_view_edit_history": true,827          "wiki": false,828          "post_url": "/t/valueerror-need-more-than-1-value-to-unpack/64697/3",829          "can_accept_answer": false,830          "can_unaccept_answer": false,831          "accepted_answer": false,832          "topic_accepted_answer": null833        },834        {835          "id": 154351,836          "name": "",837          "username": "ptrblck",838          "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png",839          "created_at": "2019-12-22T08:25:34.657Z",840          "cooked": "<aside class=\"quote no-group\" data-username=\"fupanbo\" data-post=\"3\" data-topic=\"64697\">\n<div class=\"title\">\n<div class=\"quote-controls\"></div>\n<img loading=\"lazy\" alt=\"\" width=\"24\" height=\"24\" src=\"https://discuss.pytorch.org/letter_avatar_proxy/v4/letter/f/ecae2f/48.png\" class=\"avatar\"> fupanbo:</div>\n<blockquote>\n<p>should i set them as the Variable</p>\n</blockquote>\n</aside>\n<p>No, <code>Variables</code> are deprecated since PyTorch <code>0.4.0</code>, so you can just use tensors in newer versions.</p>\n<aside class=\"quote no-group\" data-username=\"fupanbo\" data-post=\"3\" data-topic=\"64697\">\n<div class=\"title\">\n<div class=\"quote-controls\"></div>\n<img loading=\"lazy\" alt=\"\" width=\"24\" height=\"24\" src=\"https://discuss.pytorch.org/letter_avatar_proxy/v4/letter/f/ecae2f/48.png\" class=\"avatar\"> fupanbo:</div>\n<blockquote>\n<p>when the back propagation is running , it looks like the grad update will be effected?</p>\n</blockquote>\n</aside>\n<p><code>torch.split</code> and <code>torch.cat</code> won’t detach the tensors from the computation graph and the backward call will still work.</p>\n<aside class=\"quote no-group\" data-username=\"fupanbo\" data-post=\"3\" data-topic=\"64697\">\n<div class=\"title\">\n<div class=\"quote-controls\"></div>\n<img loading=\"lazy\" alt=\"\" width=\"24\" height=\"24\" src=\"https://discuss.pytorch.org/letter_avatar_proxy/v4/letter/f/ecae2f/48.png\" class=\"avatar\"> fupanbo:</div>\n<blockquote>\n<p>so the code for the part of \"2X24 turn to 2X8 \" should put in my network 's def forward<br>\nor in my every train epoch</p>\n</blockquote>\n</aside>\n<p>It doesn’t matter where these operations are called. As long as these ops are called somewhere, the results will be the same.<br>\nI would personally put these operations into the model and just return the output, which can be passed to a criterion, but your coding style might differ. <img src=\"https://discuss.pytorch.org/images/emoji/apple/wink.png?v=12\" title=\":wink:\" class=\"emoji\" alt=\":wink:\" loading=\"lazy\" width=\"20\" height=\"20\"></p>",841          "post_number": 4,842          "post_type": 1,843          "posts_count": 10,844          "updated_at": "2019-12-22T08:25:34.657Z",845          "reply_count": 1,846          "reply_to_post_number": 3,847          "quote_count": 1,848          "incoming_link_count": 4,849          "reads": 9,850          "readers_count": 8,851          "score": 26.8,852          "yours": false,853          "topic_id": 64697,854          "topic_slug": "valueerror-need-more-than-1-value-to-unpack",855          "display_username": "",856          "primary_group_name": null,857          "flair_name": null,858          "flair_url": null,859          "flair_bg_color": null,860          "flair_color": null,861          "flair_group_id": null,862          "badges_granted": [],863          "version": 1,864          "can_edit": false,865          "can_delete": false,866          "can_recover": false,867          "can_see_hidden_post": false,868          "can_wiki": false,869          "read": true,870          "user_title": "",871          "bookmarked": false,872          "actions_summary": [],873          "moderator": true,874          "admin": true,875          "staff": true,876          "user_id": 3534,877          "hidden": false,878          "trust_level": 2,879          "deleted_at": null,880          "user_deleted": false,881          "edit_reason": null,882          "can_view_edit_history": true,883          "wiki": false,884          "post_url": "/t/valueerror-need-more-than-1-value-to-unpack/64697/4",885          "can_accept_answer": false,886          "can_unaccept_answer": false,887          "accepted_answer": false,888          "topic_accepted_answer": null889        },890        {891          "id": 154352,892          "name": "Fupanbo",893          "username": "fupanbo",894          "avatar_template": "/letter_avatar_proxy/v4/letter/f/ecae2f/{size}.png",895          "created_at": "2019-12-22T08:48:54.213Z",896          "cooked": "<aside class=\"quote no-group\" data-username=\"ptrblck\" data-post=\"4\" data-topic=\"64697\">\n<div class=\"title\">\n<div class=\"quote-controls\"></div>\n<img loading=\"lazy\" alt=\"\" width=\"24\" height=\"24\" src=\"https://discuss.pytorch.org/user_avatar/discuss.pytorch.org/ptrblck/48/1823_2.png\" class=\"avatar\"> ptrblck:</div>\n<blockquote>\n<p>No, <code>Variables</code> are deprecated since PyTorch <code>0.4.0</code> , so you can just use tensors in newer versions.</p>\n</blockquote>\n</aside>\n<p>yeah i know  and my study code is based on 0.3.0 now i have not much time to change the version:pleading_face:<br>\ni want to know  outside \"the for loop \"   should the  b1,b2 need be  defined as tensor such as:b1=torch.randn(1,24)</p>\n<p>or  just as a local value  in “the for loop”</p>\n<p>other question i already know what u mean   so appreciated:kissing_heart:</p>",897          "post_number": 5,898          "post_type": 1,899          "posts_count": 10,900          "updated_at": "2019-12-22T08:48:54.213Z",901          "reply_count": 1,902          "reply_to_post_number": 4,903          "quote_count": 1,904          "incoming_link_count": 0,905          "reads": 7,906          "readers_count": 6,907          "score": 6.4,908          "yours": false,909          "topic_id": 64697,910          "topic_slug": "valueerror-need-more-than-1-value-to-unpack",911          "display_username": "Fupanbo",912          "primary_group_name": null,913          "flair_name": null,914          "flair_url": null,915          "flair_bg_color": null,916          "flair_color": null,917          "flair_group_id": null,918          "badges_granted": [],919          "version": 1,920          "can_edit": false,921          "can_delete": false,922          "can_recover": false,923          "can_see_hidden_post": false,924          "can_wiki": false,925          "read": true,926          "user_title": null,927          "bookmarked": false,928          "actions_summary": [],929          "moderator": false,930          "admin": false,931          "staff": false,932          "user_id": 25802,933          "hidden": false,934          "trust_level": 1,935          "deleted_at": null,936          "user_deleted": false,937          "edit_reason": null,938          "can_view_edit_history": true,939          "wiki": false,940          "post_url": "/t/valueerror-need-more-than-1-value-to-unpack/64697/5",941          "can_accept_answer": false,942          "can_unaccept_answer": false,943          "accepted_answer": false,944          "topic_accepted_answer": null945        },946        {947          "id": 154422,948          "name": "",949          "username": "ptrblck",950          "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png",951          "created_at": "2019-12-22T23:33:37.921Z",952          "cooked": "<p>You don’t need to define <code>b1</code> and <code>b2</code> before the for loop, if you create them via the <code>split</code> inside the loop.</p>",953          "post_number": 6,954          "post_type": 1,955          "posts_count": 10,956          "updated_at": "2019-12-23T02:59:34.669Z",957          "reply_count": 1,958          "reply_to_post_number": 5,959          "quote_count": 0,960          "incoming_link_count": 0,961          "reads": 6,962          "readers_count": 5,963          "score": 6.2,964          "yours": false,965          "topic_id": 64697,966          "topic_slug": "valueerror-need-more-than-1-value-to-unpack",967          "display_username": "",968          "primary_group_name": null,969          "flair_name": null,970          "flair_url": null,971          "flair_bg_color": null,972          "flair_color": null,973          "flair_group_id": null,974          "badges_granted": [],975          "version": 1,976          "can_edit": false,977          "can_delete": false,978          "can_recover": false,979          "can_see_hidden_post": false,980          "can_wiki": false,981          "read": true,982          "user_title": "",983          "reply_to_user": {984            "id": 25802,985            "username": "fupanbo",986            "name": "Fupanbo",987            "avatar_template": "/letter_avatar_proxy/v4/letter/f/ecae2f/{size}.png"988          },989          "bookmarked": false,990          "actions_summary": [],991          "moderator": true,992          "admin": true,993          "staff": true,994          "user_id": 3534,995          "hidden": false,996          "trust_level": 2,997          "deleted_at": null,998          "user_deleted": false,999          "edit_reason": null,1000          "can_view_edit_history": true,1001          "wiki": false,1002          "post_url": "/t/valueerror-need-more-than-1-value-to-unpack/64697/6",1003          "can_accept_answer": false,1004          "can_unaccept_answer": false,1005          "accepted_answer": false,1006          "topic_accepted_answer": null1007        },1008        {1009          "id": 154437,1010          "name": "Fupanbo",1011          "username": "fupanbo",1012          "avatar_template": "/letter_avatar_proxy/v4/letter/f/ecae2f/{size}.png",1013          "created_at": "2019-12-23T02:59:55.293Z",1014          "cooked": "<p>thank u<br>\ni recently read some literature their code is based on pytorch 0.4.0 or 0.4.1<br>\ndo u have some suggestions about the chioce between pytorch 0.4 and 1.0 or newer versions?</p>",1015          "post_number": 8,1016          "post_type": 1,1017          "posts_count": 10,1018          "updated_at": "2019-12-23T02:59:55.293Z",1019          "reply_count": 1,1020          "reply_to_post_number": 6,1021          "quote_count": 0,1022          "incoming_link_count": 1,1023          "reads": 5,1024          "readers_count": 4,1025          "score": 11.0,1026          "yours": false,1027          "topic_id": 64697,1028          "topic_slug": "valueerror-need-more-than-1-value-to-unpack",1029          "display_username": "Fupanbo",1030          "primary_group_name": null,1031          "flair_name": null,1032          "flair_url": null,1033          "flair_bg_color": null,1034          "flair_color": null,1035          "flair_group_id": null,1036          "badges_granted": [],1037          "version": 1,1038          "can_edit": false,1039          "can_delete": false,1040          "can_recover": false,1041          "can_see_hidden_post": false,1042          "can_wiki": false,1043          "read": true,1044          "user_title": null,1045          "reply_to_user": {1046            "id": 3534,1047            "username": "ptrblck",1048            "name": "",1049            "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png"1050          },1051          "bookmarked": false,1052          "actions_summary": [],1053          "moderator": false,1054          "admin": false,1055          "staff": false,1056          "user_id": 25802,1057          "hidden": false,1058          "trust_level": 1,1059          "deleted_at": null,1060          "user_deleted": false,1061          "edit_reason": null,1062          "can_view_edit_history": true,1063          "wiki": false,1064          "post_url": "/t/valueerror-need-more-than-1-value-to-unpack/64697/8",1065          "can_accept_answer": false,1066          "can_unaccept_answer": false,1067          "accepted_answer": false,1068          "topic_accepted_answer": null1069        },1070        {1071          "id": 154448,1072          "name": "",1073          "username": "ptrblck",1074          "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png",1075          "created_at": "2019-12-23T05:56:33.766Z",1076          "cooked": "<p>I would strongly recommend to use the latest stable release, which is <code>1.3.1</code> at the moment.<br>\nOtherwise you might run into legacy issues, which were already solved. Also, you’ll get all new features.</p>",1077          "post_number": 9,1078          "post_type": 1,1079          "posts_count": 10,1080          "updated_at": "2019-12-23T05:56:33.766Z",1081          "reply_count": 1,1082          "reply_to_post_number": 8,1083          "quote_count": 0,1084          "incoming_link_count": 1,1085          "reads": 5,1086          "readers_count": 4,1087          "score": 11.0,1088          "yours": false,1089          "topic_id": 64697,1090          "topic_slug": "valueerror-need-more-than-1-value-to-unpack",1091          "display_username": "",1092          "primary_group_name": null,1093          "flair_name": null,1094          "flair_url": null,1095          "flair_bg_color": null,1096          "flair_color": null,1097          "flair_group_id": null,1098          "badges_granted": [],1099          "version": 1,1100          "can_edit": false,1101          "can_delete": false,1102          "can_recover": false,1103          "can_see_hidden_post": false,1104          "can_wiki": false,1105          "read": true,1106          "user_title": "",1107          "reply_to_user": {1108            "id": 25802,1109            "username": "fupanbo",1110            "name": "Fupanbo",1111            "avatar_template": "/letter_avatar_proxy/v4/letter/f/ecae2f/{size}.png"1112          },1113          "bookmarked": false,1114          "actions_summary": [],1115          "moderator": true,1116          "admin": true,1117          "staff": true,1118          "user_id": 3534,1119          "hidden": false,1120          "trust_level": 2,1121          "deleted_at": null,1122          "user_deleted": false,1123          "edit_reason": null,1124          "can_view_edit_history": true,1125          "wiki": false,1126          "post_url": "/t/valueerror-need-more-than-1-value-to-unpack/64697/9",1127          "can_accept_answer": false,1128          "can_unaccept_answer": false,1129          "accepted_answer": false,1130          "topic_accepted_answer": null1131        },1132        {1133          "id": 154465,1134          "name": "Fupanbo",1135          "username": "fupanbo",1136          "avatar_template": "/letter_avatar_proxy/v4/letter/f/ecae2f/{size}.png",1137          "created_at": "2019-12-23T06:35:57.269Z",1138          "cooked": "<p>ok   as a college student now i get ready for  my new trip to e study pytorch<br>\nthank u for your help and patience all the time~</p>",1139          "post_number": 10,1140          "post_type": 1,1141          "posts_count": 10,1142          "updated_at": "2019-12-23T06:35:57.269Z",1143          "reply_count": 1,1144          "reply_to_post_number": 9,1145          "quote_count": 0,1146          "incoming_link_count": 7,1147          "reads": 5,1148          "readers_count": 4,1149          "score": 41.0,1150          "yours": false,1151          "topic_id": 64697,1152          "topic_slug": "valueerror-need-more-than-1-value-to-unpack",1153          "display_username": "Fupanbo",1154          "primary_group_name": null,1155          "flair_name": null,1156          "flair_url": null,1157          "flair_bg_color": null,1158          "flair_color": null,1159          "flair_group_id": null,1160          "badges_granted": [],1161          "version": 1,1162          "can_edit": false,1163          "can_delete": false,1164          "can_recover": false,1165          "can_see_hidden_post": false,1166          "can_wiki": false,1167          "read": true,1168          "user_title": null,1169          "reply_to_user": {1170            "id": 3534,1171            "username": "ptrblck",1172            "name": "",1173            "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png"1174          },1175          "bookmarked": false,1176          "actions_summary": [],1177          "moderator": false,1178          "admin": false,1179          "staff": false,1180          "user_id": 25802,1181          "hidden": false,1182          "trust_level": 1,1183          "deleted_at": null,1184          "user_deleted": false,1185          "edit_reason": null,1186          "can_view_edit_history": true,1187          "wiki": false,1188          "post_url": "/t/valueerror-need-more-than-1-value-to-unpack/64697/10",1189          "can_accept_answer": false,1190          "can_unaccept_answer": false,1191          "accepted_answer": false,1192          "topic_accepted_answer": null1193        },1194        {1195          "id": 154468,1196          "name": "",1197          "username": "ptrblck",1198          "avatar_template": "/user_avatar/discuss.pytorch.org/ptrblck/{size}/1823_2.png",1199          "created_at": "2019-12-23T06:40:14.240Z",1200          "cooked": "<p>Sure! Feel free to post your questions in this board (or search for related questions), in case you get stuck. <img src=\"https://discuss.pytorch.org/images/emoji/apple/wink.png?v=9\" title=\":wink:\" class=\"emoji\" alt=\":wink:\"></p>",

Showing the first 1,200 of 61955 lines. Download the file for the rest.