Update BM25S model
Browse files- README.md +3 -3
- corpus.jsonl +2 -2
- corpus.mmindex.json +0 -0
- data.csc.index.npy +2 -2
- indices.csc.index.npy +2 -2
- indptr.csc.index.npy +2 -2
- params.index.json +1 -1
- vocab.index.json +2 -2
README.md
CHANGED
@@ -123,9 +123,9 @@ This dataset was created using the following data:
|
|
123 |
|
124 |
| Statistic | Value |
|
125 |
| --- | --- |
|
126 |
-
| Number of documents |
|
127 |
-
| Number of tokens |
|
128 |
-
| Average tokens per document |
|
129 |
|
130 |
## Parameters
|
131 |
|
|
|
123 |
|
124 |
| Statistic | Value |
|
125 |
| --- | --- |
|
126 |
+
| Number of documents | 347536 |
|
127 |
+
| Number of tokens | 2831191 |
|
128 |
+
| Average tokens per document | 8.15 |
|
129 |
|
130 |
## Parameters
|
131 |
|
corpus.jsonl
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
-
size
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:70cbea160b6e2b451f7895ef633be9f65f8f277a4cf719f2d6d1307e07b18944
|
3 |
+
size 29870605
|
corpus.mmindex.json
CHANGED
The diff for this file is too large to render.
See raw diff
|
|
data.csc.index.npy
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
-
size
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:10f3d0fd57448dd4dcd62cf39dfc704a7e351a384b3dfce6c7d592c598194e84
|
3 |
+
size 11324892
|
indices.csc.index.npy
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
-
size
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:2fd3bbbecb5cb9da6f930c947e99d143de87670d8f06f2929fbf217d26703450
|
3 |
+
size 11324892
|
indptr.csc.index.npy
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
-
size
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:52277c0ab954d867d9d502276cc4236c6a4ffa6e81d4af199fd26a4a686589fa
|
3 |
+
size 577132
|
params.index.json
CHANGED
@@ -6,7 +6,7 @@
|
|
6 |
"idf_method": "lucene",
|
7 |
"dtype": "float32",
|
8 |
"int_dtype": "int32",
|
9 |
-
"num_docs":
|
10 |
"version": "0.2.6",
|
11 |
"backend": "numpy"
|
12 |
}
|
|
|
6 |
"idf_method": "lucene",
|
7 |
"dtype": "float32",
|
8 |
"int_dtype": "int32",
|
9 |
+
"num_docs": 347536,
|
10 |
"version": "0.2.6",
|
11 |
"backend": "numpy"
|
12 |
}
|
vocab.index.json
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
-
size
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:5262817bfcf340fbc58ca8779dce06fc99093abe768db9ebabb1099bbccb0774
|
3 |
+
size 2436979
|