-
Notifications
You must be signed in to change notification settings - Fork 14
Expand file tree
/
Copy pathsort.py
More file actions
114 lines (95 loc) · 3.97 KB
/
Copy pathsort.py
File metadata and controls
114 lines (95 loc) · 3.97 KB
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
35
36
37
38
39
40
41
42
43
44
45
46
47
48
49
50
51
52
53
54
55
56
57
58
59
60
61
62
63
64
65
66
67
68
69
70
71
72
73
74
75
76
77
78
79
80
81
82
83
84
85
86
87
88
89
90
91
92
93
94
95
96
97
98
99
100
101
102
103
104
105
106
107
108
109
110
111
112
113
114
# Pyrogram - Telegram MTProto API Client Library for Python
# Copyright (C) 2017-present Dan <https://github.com/delivrance>
#
# This file is part of Pyrogram.
#
# Pyrogram is free software: you can redistribute it and/or modify
# it under the terms of the GNU Lesser General Public License as published
# by the Free Software Foundation, either version 3 of the License, or
# (at your option) any later version.
#
# Pyrogram is distributed in the hope that it will be useful,
# but WITHOUT ANY WARRANTY; without even the implied warranty of
# MERCHANTABILITY or FITNESS FOR A PARTICULAR PURPOSE. See the
# GNU Lesser General Public License for more details.
#
# You should have received a copy of the GNU Lesser General Public License
# along with Pyrogram. If not, see <http://www.gnu.org/licenses/>.
import csv
from pathlib import Path
import re
import requests # requests==2.28.1
import sys
if len(sys.argv) != 2:
sys.exit(1)
if sys.argv[1] == "sort":
for p in Path("source").glob("*.tsv"):
with open(p) as f:
reader = csv.reader(f, delimiter="\t")
dct = {k: v for k, v in reader if k != "id"}
keys = sorted(dct)
with open(p, "w") as f:
f.write("id\tmessage\n")
for i, item in enumerate(keys, start=1):
f.write(f"{item}\t{dct[item]}")
if i != len(keys):
f.write("\n")
elif sys.argv[1] == "scrape":
base_url = "https://core.telegram.org"
errors_api_path = "/api/errors"
errors_data = None
html_content = None
try:
response = requests.get(base_url + errors_api_path)
response.raise_for_status()
html_content = response.text
except requests.exceptions.RequestException as e:
print(f"Error fetching errors page: {e}")
match = re.search(r'<a href=\"(.*?)\">here.*?</a>', html_content)
if not match:
print("Link to errors JSON not found.")
errors_json_url = base_url + match.group(1)
try:
response = requests.get(errors_json_url)
response.raise_for_status()
errors_data = response.json()
except requests.exceptions.RequestException as e:
print(f"Error fetching errors JSON: {e}")
except ValueError:
print("Error decoding JSON response.")
error_categories = errors_data.get("errors", [])
if not error_categories:
print("No error categories found in JSON.")
source_dir = Path("source/")
source_dir.mkdir(exist_ok=True)
for error_type in error_categories:
data_dict = {}
for tsv_path in source_dir.glob(f"{error_type}*.tsv"):
with open(tsv_path, 'r') as tsv_file:
reader = csv.DictReader(tsv_file, delimiter='\t')
for row in reader:
if row['id'] != 'id':
data_dict[row['id']] = row['message']
error_details = errors_data["errors"].get(error_type, {})
descriptions = errors_data["descriptions"]
for error_code in error_details:
if error_code.endswith("_*"):
continue
description = descriptions.get(error_code, "")
if description:
processed_description = (
description.replace("%d", "{value}")
.replace("\"", "'")
.replace("»", "»")
.replace("«", "«")
.replace(" »", "")
.replace("](/api/", f"]({base_url}/api/")
)
data_dict[error_code.replace("_%d", "_X")] = processed_description
sorted_keys = sorted(data_dict.keys())
for tsv_path in source_dir.glob(f"{error_type}*.tsv"):
print(f"Writing to {tsv_path}")
with open(tsv_path, 'w', newline='') as tsv_file:
writer = csv.writer(tsv_file, delimiter='\t')
writer.writerow(["id", "message"])
writer.writerows([(key, data_dict[key]) for key in sorted_keys])