mirror of
https://github.com/bellingcat/vk-url-scraper.git
synced 2026-06-12 21:38:36 +03:00
Compare commits
10 Commits
| Author | SHA1 | Date | |
|---|---|---|---|
|
|
f67707a740 | ||
|
|
798684a334 | ||
|
|
a556b237e9 | ||
|
|
283bc35658 | ||
|
|
cef70fb80d | ||
|
|
e66ef4f477 | ||
|
|
1f6a8368fd | ||
|
|
9a046fd1cb | ||
|
|
aae2bb5999 | ||
|
|
9e30b81d16 |
1020
Pipfile.lock
generated
1020
Pipfile.lock
generated
File diff suppressed because it is too large
Load Diff
@@ -40,3 +40,4 @@ sphinx-autodoc-typehints
|
|||||||
|
|
||||||
# For parsing and comparing version numbers.
|
# For parsing and comparing version numbers.
|
||||||
packaging
|
packaging
|
||||||
|
python-dotenv==0.21.1
|
||||||
@@ -5,11 +5,15 @@
|
|||||||
# pipenv lock --requirements
|
# pipenv lock --requirements
|
||||||
#
|
#
|
||||||
|
|
||||||
certifi==2022.6.15
|
-i https://pypi.org/simple
|
||||||
charset-normalizer==2.0.12
|
brotli==1.0.9; platform_python_implementation == 'CPython'
|
||||||
idna==3.3
|
certifi==2022.12.7; python_version >= '3.6'
|
||||||
requests==2.28.0
|
charset-normalizer==3.0.1; python_version >= '3.6'
|
||||||
urllib3==1.26.9
|
idna==3.4; python_version >= '3.5'
|
||||||
vk-api==11.9.8
|
mutagen==1.46.0; python_version >= '3.7'
|
||||||
python-dotenv==0.20.0
|
pycryptodomex==3.17; python_version >= '2.7' and python_version not in '3.0, 3.1, 3.2, 3.3, 3.4'
|
||||||
yt-dlp==2022.7.18
|
requests==2.28.2; python_version >= '3.7' and python_version < '4'
|
||||||
|
urllib3==1.26.14; python_version >= '2.7' and python_version not in '3.0, 3.1, 3.2, 3.3, 3.4, 3.5'
|
||||||
|
vk-api==11.9.9
|
||||||
|
websockets==10.4; python_version >= '3.7'
|
||||||
|
yt-dlp==2023.2.17
|
||||||
5
setup.py
5
setup.py
@@ -44,7 +44,10 @@ setup(
|
|||||||
"Programming Language :: Python :: 3",
|
"Programming Language :: Python :: 3",
|
||||||
],
|
],
|
||||||
keywords=["scraper", "vk", "vkontakte", "vk-api", "media-downloader"],
|
keywords=["scraper", "vk", "vkontakte", "vk-api", "media-downloader"],
|
||||||
url="https://github.com/bellingcat/vk-url-scraper",
|
project_urls={
|
||||||
|
"Code": "https://github.com/bellingcat/vk-url-scraper",
|
||||||
|
"Documentation": "https://vk-url-scraper.readthedocs.io/en/latest/",
|
||||||
|
},
|
||||||
author="Bellingcat",
|
author="Bellingcat",
|
||||||
author_email="tech@bellingcat.com",
|
author_email="tech@bellingcat.com",
|
||||||
license="MIT",
|
license="MIT",
|
||||||
|
|||||||
@@ -19,7 +19,6 @@ def test_login_custom_file():
|
|||||||
VkScraper(
|
VkScraper(
|
||||||
os.environ["VK_USERNAME"],
|
os.environ["VK_USERNAME"],
|
||||||
os.environ["VK_PASSWORD"],
|
os.environ["VK_PASSWORD"],
|
||||||
os.environ.get("VK_TOKEN"),
|
|
||||||
session_file=session_filename,
|
session_file=session_filename,
|
||||||
)
|
)
|
||||||
assert os.path.isfile(session_filename)
|
assert os.path.isfile(session_filename)
|
||||||
@@ -81,7 +80,7 @@ def test_scrape_wall_url_with_photos():
|
|||||||
== "Хабаровск\nАллея героев\nПомолимся об укокоении воинов:\nАлександра, Игоря, Эдуарда, \nДионисия, Евгения, Александра, Артемия, Иннокентия, Андрея."
|
== "Хабаровск\nАллея героев\nПомолимся об укокоении воинов:\nАлександра, Игоря, Эдуарда, \nДионисия, Евгения, Александра, Артемия, Иннокентия, Андрея."
|
||||||
)
|
)
|
||||||
assert str(res[0]["datetime"]) == str(datetime.datetime(2022, 6, 15, 10, 37, 24))
|
assert str(res[0]["datetime"]) == str(datetime.datetime(2022, 6, 15, 10, 37, 24))
|
||||||
assert len(res[0]["payload"]) == 16
|
assert len(res[0]["payload"]) == 17
|
||||||
assert len(res[0]["attachments"].keys()) == 1
|
assert len(res[0]["attachments"].keys()) == 1
|
||||||
assert list(res[0]["attachments"].keys()) == ["photo"]
|
assert list(res[0]["attachments"].keys()) == ["photo"]
|
||||||
assert len(res[0]["attachments"]["photo"]) == 9
|
assert len(res[0]["attachments"]["photo"]) == 9
|
||||||
@@ -93,7 +92,7 @@ def test_scrape_wall_url_with_photos_inner_videos_and_links_with_inner_photos():
|
|||||||
assert res[0]["id"] == "wall-17315087_74182"
|
assert res[0]["id"] == "wall-17315087_74182"
|
||||||
assert res[0]["text"] == ""
|
assert res[0]["text"] == ""
|
||||||
assert str(res[0]["datetime"]) == str(datetime.datetime(2022, 3, 24, 11, 1, 9))
|
assert str(res[0]["datetime"]) == str(datetime.datetime(2022, 3, 24, 11, 1, 9))
|
||||||
assert len(res[0]["payload"]) == 15
|
assert len(res[0]["payload"]) == 17
|
||||||
assert len(res[0]["attachments"].keys()) == 3
|
assert len(res[0]["attachments"].keys()) == 3
|
||||||
for k in ["photo", "link", "video"]:
|
for k in ["photo", "link", "video"]:
|
||||||
assert k in list(res[0]["attachments"].keys())
|
assert k in list(res[0]["attachments"].keys())
|
||||||
|
|||||||
@@ -61,7 +61,7 @@ class VkScraper:
|
|||||||
token : str
|
token : str
|
||||||
Access token received after authenticating, can be found in the vl_config.v2.json file
|
Access token received after authenticating, can be found in the vl_config.v2.json file
|
||||||
session_file : str
|
session_file : str
|
||||||
File name where the VK session is saved so future logins are easier
|
File name where the VK session is saved so future logins are easier, this will not be created if token is passed
|
||||||
captcha_handler : func
|
captcha_handler : func
|
||||||
Function that can receive a vk_api captcha instance and help the user solve it, default is a complete CLI handler
|
Function that can receive a vk_api captcha instance and help the user solve it, default is a complete CLI handler
|
||||||
"""
|
"""
|
||||||
|
|||||||
@@ -2,7 +2,7 @@ _MAJOR = "0"
|
|||||||
_MINOR = "3"
|
_MINOR = "3"
|
||||||
# On main and in a nightly release the patch should be one ahead of the last
|
# On main and in a nightly release the patch should be one ahead of the last
|
||||||
# released build.
|
# released build.
|
||||||
_PATCH = "6"
|
_PATCH = "14"
|
||||||
# This is mainly for nightly builds which have the suffix ".dev$DATE". See
|
# This is mainly for nightly builds which have the suffix ".dev$DATE". See
|
||||||
# https://semver.org/#is-v123-a-semantic-version for the semantics.
|
# https://semver.org/#is-v123-a-semantic-version for the semantics.
|
||||||
_SUFFIX = ""
|
_SUFFIX = ""
|
||||||
|
|||||||
Reference in New Issue
Block a user