Upload optimized language model w/ WebGPU-compatible GQA (#11)
Browse files- [work in progress] Upload optimized language model w/ WebGPU-compatible GQA (2bb3f5c5ebd5d778f62f35655fe6dc3628e4f776)
- onnx/decoder_model_merged.onnx +2 -2
- onnx/decoder_model_merged_bnb4.onnx +2 -2
- onnx/decoder_model_merged_fp16.onnx +2 -2
- onnx/decoder_model_merged_int8.onnx +2 -2
- onnx/decoder_model_merged_q4.onnx +2 -2
- onnx/decoder_model_merged_q4f16.onnx +2 -2
- onnx/decoder_model_merged_quantized.onnx +2 -2
- onnx/decoder_model_merged_uint8.onnx +2 -2
onnx/decoder_model_merged.onnx
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
-
size
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:5ae6a66fdf47680433aa6bcc512ae8f0b98c2db01bea055fc76ee37bb0389b78
|
3 |
+
size 540640302
|
onnx/decoder_model_merged_bnb4.onnx
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
-
size
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:1016b8e6ecb95aa36da3f6ab32c8c955f83297ec9f219d08009eea33e6d1089e
|
3 |
+
size 78154897
|
onnx/decoder_model_merged_fp16.onnx
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
-
size
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:925ffaf3af9336b4b4c064d77fa64128868d8cb915c257520708d3aa853bb4ae
|
3 |
+
size 270414183
|
onnx/decoder_model_merged_int8.onnx
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
-
size
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:306fad6cec3b5b7613600b74abf1e410fa50000fbc81292c0d083bd8d85b5fe2
|
3 |
+
size 137221530
|
onnx/decoder_model_merged_q4.onnx
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
-
size
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:e3fd81130e30715580cf148398e48e2925b6219bff75e17d2a77add35a363f65
|
3 |
+
size 86562901
|
onnx/decoder_model_merged_q4f16.onnx
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
-
size
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:0a6b4e72e26025c190f37e2573ea2ce3336d95670da9c4b1fa662e47cc3d94a7
|
3 |
+
size 77034560
|
onnx/decoder_model_merged_quantized.onnx
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
-
size
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:33f14f3bca52699d733d86fcc9c5a0ec6f57afffa1dc9596abde53cde9df81aa
|
3 |
+
size 137221644
|
onnx/decoder_model_merged_uint8.onnx
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
-
size
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:33f14f3bca52699d733d86fcc9c5a0ec6f57afffa1dc9596abde53cde9df81aa
|
3 |
+
size 137221644
|