[youtube] Relax URL expansion in description

[youtube-dl] / youtube_dl / extractor / vk.py
diff --git a/youtube_dl/extractor/vk.py b/youtube_dl/extractor/vk.py

index 79c819bc389a09b2cd5186f64feb49d3e474d189..cfc5ffd8b02082d1f6d3035ba7b02dc230b8b2ed 100644 (file)
--- a/youtube_dl/extractor/vk.py
+++ b/youtube_dl/extractor/vk.py
@@ -3,6 +3,7 @@ from __future__ import unicode_literals
  
  import re
  import json
+import sys
  
  from .common import InfoExtractor
  from ..compat import compat_str
@@ -10,7 +11,6 @@ from ..utils import (
      ExtractorError,
      int_or_none,
      orderedSet,
-    sanitized_Request,
      str_to_int,
      unescapeHTML,
      unified_strdate,
@@ -190,7 +190,7 @@ class VKIE(InfoExtractor):
          if username is None:
              return
  
-        login_page = self._download_webpage(
+        login_page, url_handle = self._download_webpage_handle(
              'https://vk.com', None, 'Downloading login page')
  
          login_form = self._hidden_inputs(login_page)
@@ -200,11 +200,26 @@ class VKIE(InfoExtractor):
              'pass': password.encode('cp1251'),
          })
  
-        request = sanitized_Request(
-            'https://login.vk.com/?act=login',
-            urlencode_postdata(login_form))
+        # https://new.vk.com/ serves two same remixlhk cookies in Set-Cookie header
+        # and expects the first one to be set rather than second (see
+        # https://github.com/rg3/youtube-dl/issues/9841#issuecomment-227871201).
+        # As of RFC6265 the newer one cookie should be set into cookie store
+        # what actually happens.
+        # We will workaround this VK issue by resetting the remixlhk cookie to
+        # the first one manually.
+        cookies = url_handle.headers.get('Set-Cookie')
+        if sys.version_info[0] >= 3:
+            cookies = cookies.encode('iso-8859-1')
+        cookies = cookies.decode('utf-8')
+        remixlhk = re.search(r'remixlhk=(.+?);.*?\bdomain=(.+?)(?:[,;]|$)', cookies)
+        if remixlhk:
+            value, domain = remixlhk.groups()
+            self._set_cookie(domain, 'remixlhk', value)
+
          login_page = self._download_webpage(
-            request, None, note='Logging in as %s' % username)
+            'https://login.vk.com/?act=login', None,
+            note='Logging in as %s' % username,
+            data=urlencode_postdata(login_form))
  
          if re.search(r'onLoginFailed', login_page):
              raise ExtractorError(