diff --git a/model-00001-of-000163.safetensors b/model-00001-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..5ab5b8e923ca36cdfc43356c6bb5d326b24f71de --- /dev/null +++ b/model-00001-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:e5e87b89ead036e72853e1d7cf48136d40d85b90cc242944e58aa7f88900092f +size 8609454256 diff --git a/model-00002-of-000163.safetensors b/model-00002-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..3723cdc6ac74103e2de31c5a20bb452cb6938929 --- /dev/null +++ b/model-00002-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:b620aa86006f54aac8d8453e9370cb8411119e3038e13df0bce82ebd954729f1 +size 8602553952 diff --git a/model-00003-of-000163.safetensors b/model-00003-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..f477d89a00072d5e1d6d3feee7aee5bf05cc144c --- /dev/null +++ b/model-00003-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:0be9a653c37d1b3d0817a4721349491cee42c9e5ff0c8658a31dc2291cf1db14 +size 8602554152 diff --git a/model-00004-of-000163.safetensors b/model-00004-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..dda9a76e7d51bc68db866d32b6ba9df8fca74f72 --- /dev/null +++ b/model-00004-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:3a9fffda0b2c62ba57e87a9a246113fc3ceb45aa256cf9ff9fdd75dd345a2521 +size 8598786296 diff --git a/model-00005-of-000163.safetensors b/model-00005-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..9c7bdc7a83d0c2997a922fe762de325b89750049 --- /dev/null +++ b/model-00005-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:59baf5fe1862df4db604533931b07fcb97596712a094851cf85793dc4452c268 +size 8602554048 diff --git a/model-00006-of-000163.safetensors b/model-00006-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..580bee5c7888356a962f6f43b6404dc832203f09 --- /dev/null +++ b/model-00006-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:2cfa5d5f02c7ab85d47e07e82cb9353232a842aad2f4747b01489956bf4845df +size 8741916520 diff --git a/model-00007-of-000163.safetensors b/model-00007-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..f2e037a7d7b59fbf9eb4146d9108892a65cbf026 --- /dev/null +++ b/model-00007-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:4b6f2cdd5392caaa845e140b7f178e4406043771f0d277ad1d92b5cf36a3238b +size 8606225096 diff --git a/model-00008-of-000163.safetensors b/model-00008-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..e23202075aa54e6c6a901ac4b1c518fc5fdb8287 --- /dev/null +++ b/model-00008-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:fdf1644b293adcf07b9d708f5d385f43575fd2fe962c6d766e48c84f1684d235 +size 8602554144 diff --git a/model-00009-of-000163.safetensors b/model-00009-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..42f5572a5f437cbcd3d000b544c8d23a7fea2d4c --- /dev/null +++ b/model-00009-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:b36db93b74136b8852d062c85e764ae4356c3857751c5430d3c33e01bf2fbf91 +size 8598786392 diff --git a/model-00010-of-000163.safetensors b/model-00010-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..5abc24a9eff3dc8dc69c192480b79f14fb4477d2 --- /dev/null +++ b/model-00010-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:ff5bdffc3fca98f7c43b91b12bd0df70661b656050f0fd3ac7f7f1849d580b1a +size 8602553952 diff --git a/model-00011-of-000163.safetensors b/model-00011-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..25274f247e6fa94ef5ed8dff617f581a36a471e7 --- /dev/null +++ b/model-00011-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:10faa73df07efcc7961c547b5cb93cf17d612d32ef52e9830b4260d0ce63d95a +size 8602554152 diff --git a/model-00012-of-000163.safetensors b/model-00012-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..071aab62b8ff764eb51a3bf133fa992f1b2b9ef8 --- /dev/null +++ b/model-00012-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:e204197e2af834ebfd4af43526914503fadfa58796b4705579314f4e9887324d +size 2642451624 diff --git a/model-00013-of-000163.safetensors b/model-00013-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..9d7b6572477f3d35f0b490c34d48e409f3993821 --- /dev/null +++ b/model-00013-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:d039dcb1f0f65987e36d045006f139b8d6c99fe8a9373a70c70991213cfa9f43 +size 8598757320 diff --git a/model-00014-of-000163.safetensors b/model-00014-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..a538006474d018b579e143b2dbb9c68667a6d736 --- /dev/null +++ b/model-00014-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:97d7725d9ef598faba4f1284cab74b1ae32cdfd02b828079b904a52eba0d7b8e +size 8602554136 diff --git a/model-00015-of-000163.safetensors b/model-00015-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..85d1512cba65b0ffd4e269d950c84be260c7fd16 --- /dev/null +++ b/model-00015-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:a46d3ea85ba43a66193d0f7546ec1c936a4b211fad91188ab8f43093b60f4d37 +size 8598786408 diff --git a/model-00016-of-000163.safetensors b/model-00016-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..1fe5502de8b035b46ed089b714e4bc8539a21ace --- /dev/null +++ b/model-00016-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:b6630ebc4caf694bb3918918a0a06dd00f3228c764f0725bebacc7d7c02c935e +size 8602553936 diff --git a/model-00017-of-000163.safetensors b/model-00017-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..8469c5cc489ee8ab4323d167115d5f1677f878f1 --- /dev/null +++ b/model-00017-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:becb73788e18f5b3ab21e4aaa0ed846791b9b062c247fed8daf48da3006e463c +size 8602554152 diff --git a/model-00018-of-000163.safetensors b/model-00018-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..f7bc4b2d46db79f58171d5a6ad982e2232d2366a --- /dev/null +++ b/model-00018-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:3b9d024b0e6b87889bdae11855048dd7dc26df21bc1b354962b270cfb7ea551c +size 8598786312 diff --git a/model-00019-of-000163.safetensors b/model-00019-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..b13e60bf2ac9811e46e42732c52eb35ef71711a7 --- /dev/null +++ b/model-00019-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:a49fb1889fd754eb3db9b25dbea6ebb5f152b1f98b9374157306c258b997e540 +size 8602554032 diff --git a/model-00020-of-000163.safetensors b/model-00020-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..8e6d79147029ae5744884ea8788f5df099dfa554 --- /dev/null +++ b/model-00020-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:e039712e7b55545d7a5e26d99c2c37f23d354b09c773251da9d95d2f91f79ae2 +size 8602554160 diff --git a/model-00021-of-000163.safetensors b/model-00021-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..7eaa34f5188487a6452cd6bfa5b53ad8aa388dea --- /dev/null +++ b/model-00021-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:6759dceaf8e72b62eca652da8cc923102f5d00d6b2a4b17c8b2b7d582fcab5a5 +size 8598786512 diff --git a/model-00023-of-000163.safetensors b/model-00023-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..05579c1a3b6676bdeed07fa139ac9d3801ee2ccd --- /dev/null +++ b/model-00023-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:b41bc4167b28a28ebf96c29ab7d744f46398030c0de383a091c27d2b378fabb7 +size 8598786704 diff --git a/model-00024-of-000163.safetensors b/model-00024-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..8f42da3d72181ef2125e5c80a3c6eca86e683cad --- /dev/null +++ b/model-00024-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:8a1e67fa44a09a698e070f0f81f637c1bca664ed40e45f24c9a307c076ae64fc +size 8602554224 diff --git a/model-00025-of-000163.safetensors b/model-00025-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..1f4a127f9ae653753fea001bba50f4afdc8ceaeb --- /dev/null +++ b/model-00025-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:40f972cced4deee0e61071ed94ce998d099083cf4be8b431efb06ab5d105e313 +size 8602554448 diff --git a/model-00026-of-000163.safetensors b/model-00026-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..b43bf34a274f70402e9431b5b7b74722ef251246 --- /dev/null +++ b/model-00026-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:4c7bf2562f107c268579a8cdae8e87e165b0c780fd7166094d4d5632b5b04b10 +size 8598786616 diff --git a/model-00027-of-000163.safetensors b/model-00027-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..b173190acd0abc7be2c9377d228937c47a983a7d --- /dev/null +++ b/model-00027-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:c72f7f985dc12fa4e9105c55eecb9be6139b78ed17713e70013af7f375d9c929 +size 8602554312 diff --git a/model-00028-of-000163.safetensors b/model-00028-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..1e203d0e7d4b48d4b665aaef5e7eba415f0151bb --- /dev/null +++ b/model-00028-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:ea9febe137ad6c326eea88e093218e28fb22fea439319025eb89e5032af7f851 +size 8602554448 diff --git a/model-00029-of-000163.safetensors b/model-00029-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..21564b6062a526deac44314bc176d3ad18368e61 --- /dev/null +++ b/model-00029-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:51e0500521bba5dbd8faf94bc68af7b3b6197b67498cc5f078e4e2199b14a890 +size 8598786520 diff --git a/model-00030-of-000163.safetensors b/model-00030-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..0734f2a853a54739403af5aa17f9baa8bfd236bb --- /dev/null +++ b/model-00030-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:fc3a398e0e0dd6995c4731f7333830ae359882a994e1790d2045e1ff4e7681ed +size 8602554408 diff --git a/model-00031-of-000163.safetensors b/model-00031-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..d62effe197f2a2f31e938aa0456b2fe38ab43ef1 --- /dev/null +++ b/model-00031-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:3d81fe26a786ed8d4238b0ffd976ead048075cfd4ee5222c02b34e5727d3b373 +size 8598786720 diff --git a/model-00032-of-000163.safetensors b/model-00032-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..076500e86f61190d80907a4caa6d57469e8c8ef4 --- /dev/null +++ b/model-00032-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:e8e98fbd5a8c8ec709c86608faca725309041203a71b1ea3ee20ccac1575013a +size 8602554208 diff --git a/model-00033-of-000163.safetensors b/model-00033-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..32db6941a1f79dddf0eca38efcab694cc5af8d15 --- /dev/null +++ b/model-00033-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:16043317a49e745e5743194920d6d5cf73472d5e7581b5197407465903525196 +size 8602554448 diff --git a/model-00034-of-000163.safetensors b/model-00034-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..d5c834a9e1b6ba151e599342dd70db4cf3375bb1 --- /dev/null +++ b/model-00034-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:8bdf4e8779f20854d43608277dec089d72b547c5636c5cf6e1b1cc8a3488e18d +size 3493899088 diff --git a/model-00035-of-000163.safetensors b/model-00035-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..72cbbdfa3b9252e1397f9148aee88766ce97aa51 --- /dev/null +++ b/model-00035-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:0dc868c6f6dee0fa07b69700ff7c2669d32bbe44b4c3e809d202930f5851bbef +size 8598757608 diff --git a/model-00036-of-000163.safetensors b/model-00036-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..5e9e45e453200240048196f12a95b19c8d49434c --- /dev/null +++ b/model-00036-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:663c4175d3c9dcda7e3233e497a2b8eba354f9b5ba08f41201109e5cb111fa11 +size 8602554424 diff --git a/model-00037-of-000163.safetensors b/model-00037-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..abb0584fe3ece056b6360d797bef17d9e4bdcd83 --- /dev/null +++ b/model-00037-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:7116de8fdf43da16827efb71433a60424358a36573d060626a6e647d571326e2 +size 8598786704 diff --git a/model-00038-of-000163.safetensors b/model-00038-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..be8838dca7fb052e5fe0a8be64969df99f36ff69 --- /dev/null +++ b/model-00038-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:21723ffa0998ec533df4021b821518668f298f644d1a26b300cd8552128d6b99 +size 8602554224 diff --git a/model-00039-of-000163.safetensors b/model-00039-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..35433ae17f0fc3aca94161ec18d4de9868447f08 --- /dev/null +++ b/model-00039-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:a16ede8ff7d1e2fb96a9ade240f30aebdf9aec0c033eeab66869be3ffde9be4e +size 8602554448 diff --git a/model-00041-of-000163.safetensors b/model-00041-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..9d1304cd6817115a75abb696ad139adc9842107f --- /dev/null +++ b/model-00041-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:e77668a49e7785b1d8a5e2ceb4875c87cfe9ecac51e1ee420bc41ef8fb167a97 +size 8602554320 diff --git a/model-00042-of-000163.safetensors b/model-00042-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..44bc2e278f36bfd20eac7e0bb201836af4f4c1c6 --- /dev/null +++ b/model-00042-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:a3a0f94bdb238df8e9504fb349634fb76be37f60952ee3f64b6c529445a45766 +size 8602554448 diff --git a/model-00043-of-000163.safetensors b/model-00043-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..ba83c5dfc47b86e96833e8a4c6ae4f2cdc797968 --- /dev/null +++ b/model-00043-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:e495f6da3338565b81305ca1e2ed925f6de504c4928cd17ae7eb03ea0b8be695 +size 8598786504 diff --git a/model-00044-of-000163.safetensors b/model-00044-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..bcae9244a0d5d5eb55d1ee88aedf7e902958d298 --- /dev/null +++ b/model-00044-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:6e7947a7ac8b51f9c73731f5de9b2398238cee5ef8a9bfded3797473b498f3d0 +size 8602554416 diff --git a/model-00045-of-000163.safetensors b/model-00045-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..e222d8943b66e2db1c7bfa845cd882039230d826 --- /dev/null +++ b/model-00045-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:867c286983f1633248ad628074d591d1a75191c45786cc0610db447a5a48cdb7 +size 8598786704 diff --git a/model-00046-of-000163.safetensors b/model-00046-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..4000d42e6d79bffb7d5129753bf6f6e195face22 --- /dev/null +++ b/model-00046-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:9d1446557a2e5f6c93361788626ce0daf918c41c771125b75a08ea5a1953daca +size 8602554224 diff --git a/model-00047-of-000163.safetensors b/model-00047-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..77565ab4fee4685421823313d717c4693727495e --- /dev/null +++ b/model-00047-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:bca76a65e27b91cc7f5eda783b6c8c9f064a5d260bf2747ef025d7df08e4f873 +size 8602554448 diff --git a/model-00048-of-000163.safetensors b/model-00048-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..c82826e13bae1bb3a03c2b1f293df3573a900ea7 --- /dev/null +++ b/model-00048-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:6a484c57aa9273db48ebf95f2d13ca8548188c1c8380abdd1811fe6a3d893177 +size 8598786616 diff --git a/model-00049-of-000163.safetensors b/model-00049-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..1330f673bd42db69fee8f3524afb146862f33102 --- /dev/null +++ b/model-00049-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:58532be8d29009451a3bb3e1bc5e29236c6e80f9650595a5dea6b12784ee0dd0 +size 8602554312 diff --git a/model-00050-of-000163.safetensors b/model-00050-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..855d80e67b2a1799462a19d1a7f4b22db5291bf8 --- /dev/null +++ b/model-00050-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:59b4009aedde11ca72dbf8763aa48523ee67bce482688c4650987bbda8a781af +size 8602554448 diff --git a/model-00051-of-000163.safetensors b/model-00051-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..56c65b92a5f6f2528917111b7c2f7c79fa19cdb5 --- /dev/null +++ b/model-00051-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:484b03be0ebf2fc12fb97ae4382357b57ec43e80a17622c96f0c66abc2ec0ccc +size 8598786520 diff --git a/model-00052-of-000163.safetensors b/model-00052-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..30dc65c25fc4e80805d65262a579e4ef42edf194 --- /dev/null +++ b/model-00052-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:3fb5da4bf9365c4df89af0a867768ef16f6b1f1ba46594f04037085b911e7e5d +size 8602554408 diff --git a/model-00053-of-000163.safetensors b/model-00053-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..e1c69641d915b1646007e1f46512e248b99b09c7 --- /dev/null +++ b/model-00053-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:b60d4386d43ff3ea7852c85f9926da89950eeec8dbe28ad511063ca276739ea9 +size 8598786720 diff --git a/model-00054-of-000163.safetensors b/model-00054-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..50e751d65a3b0fb283dc58d8205a31ff91ea9b19 --- /dev/null +++ b/model-00054-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:41d8805f5b2a0b45de2a7f1ce41a4d68d0a157fd811e4d8f09f4f73c810ea7b9 +size 8602554208 diff --git a/model-00055-of-000163.safetensors b/model-00055-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..fd948a82428b901051de8fac9b50e05b4a840014 --- /dev/null +++ b/model-00055-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:fbe0e400217d2d80ca980433d6684230ca795fec3425fe0edc982f694af7a5e0 +size 8602554448 diff --git a/model-00056-of-000163.safetensors b/model-00056-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..8a286b7b0b8a27ca814264c7caf2cdb13087dd3c --- /dev/null +++ b/model-00056-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:4fdd86f08f8ffafb7236c38e04775088378bea518af4d063ac07a3f43924c1ce +size 3493899088 diff --git a/model-00057-of-000163.safetensors b/model-00057-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..9c5c5f156ec2810a4e395894e6e24700e93be29c --- /dev/null +++ b/model-00057-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:9fd1d954444702f042c6949b77d4e198ba2c8c415db5baf76cff9fcbd55895d4 +size 8598757608 diff --git a/model-00058-of-000163.safetensors b/model-00058-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..c9466faa4a34bc3b3f200a12d43c141b206422cd --- /dev/null +++ b/model-00058-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:642cd14a735600af3d00c9b9187e9aaf2ad6b9174681193ddf270fe89ee2bfaa +size 8602554424 diff --git a/model-00059-of-000163.safetensors b/model-00059-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..451ec54f5401deed32caba79c3d144112b4ed96d --- /dev/null +++ b/model-00059-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:0a8543b0ccc8cb73df1044e80b78d74beba2c233e2cca05989f0fd7e7e0fa42b +size 8598786704 diff --git a/model-00060-of-000163.safetensors b/model-00060-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..a693654a194020471d5b2ac1678f3ad967ed6b8c --- /dev/null +++ b/model-00060-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:d4861e18bedec884d6191ae27386f4b3f2a0b52f6be57df991b1b9d0ad751fa9 +size 8602554224 diff --git a/model-00061-of-000163.safetensors b/model-00061-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..ac6d444de6babf327941f8cd92d03f128e20983c --- /dev/null +++ b/model-00061-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:ddbd126b630a8cc954077706d1469ea2df8a5c11677ce336e3d30808b312431c +size 8602554448 diff --git a/model-00064-of-000163.safetensors b/model-00064-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..b5c0ca1aa8a994a58a96295da322489fbba33e9d --- /dev/null +++ b/model-00064-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:31fa21080179721762ce4c568da7d7c49ea7aa76dc598ec7b1033da6011087fd +size 8602554448 diff --git a/model-00065-of-000163.safetensors b/model-00065-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..b67a43c44a56204ca1a1027912daa8533ac65e7f --- /dev/null +++ b/model-00065-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:7ba6c10d38d6b3781d49b5396ae907b854fcec20b7b1966c0091eccff9702d73 +size 8598786504 diff --git a/model-00066-of-000163.safetensors b/model-00066-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..081efc8d423b22b3e71f6423139a7f20a515e588 --- /dev/null +++ b/model-00066-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:83e856db43bfb8aab0c9fa752c2ba89c95d3c4daf90fbfed872da2e8e164f26f +size 8602554416 diff --git a/model-00067-of-000163.safetensors b/model-00067-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..0d8fc1db3121401fc8e5c77dafa016d2499faa93 --- /dev/null +++ b/model-00067-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:e5c723b79180ef12f2e1a82b1600d2e415dd966c9c4166ff07b46ebc1dee00fa +size 8598786704 diff --git a/model-00068-of-000163.safetensors b/model-00068-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..73f1ec02dd192e0236490c6e644baf7a016c5a95 --- /dev/null +++ b/model-00068-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:80b8e57578494d37fc163bed4a5b3b8ad03ba218a9d5f8af78610b1f8580ee47 +size 8602554224 diff --git a/model-00069-of-000163.safetensors b/model-00069-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..ffd247c61e8cc9d05db42113be39f551e43ff251 --- /dev/null +++ b/model-00069-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:1eac93e819183abf53a2bca5b5e4a39ccd7368343f28d5103b7c98cd2832574d +size 8602554448 diff --git a/model-00070-of-000163.safetensors b/model-00070-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..430042a556d54e835896bd3067494d66aba5c329 --- /dev/null +++ b/model-00070-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:ce2107a2c3b62cf21cacc86b7997ad54f4f2b2764015fbf0649039c084a97454 +size 8598786616 diff --git a/model-00071-of-000163.safetensors b/model-00071-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..1d70d47e548313c445d21bf8dcb3eee2f6417d1a --- /dev/null +++ b/model-00071-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:b141897abcb5642b33b694f889c153faea38defcbb87a672ef30dc1dc684730c +size 8602554312 diff --git a/model-00072-of-000163.safetensors b/model-00072-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..3bf11d23e7f03b7baa95ce6ccc4c0cfc335e28af --- /dev/null +++ b/model-00072-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:5d4d9ef4afb7873dfa001f71f6b44ffd2103bcf627a012d792185821be16c85e +size 8602554448 diff --git a/model-00073-of-000163.safetensors b/model-00073-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..e424a65470a13f777b5eafca82ac32a4269ba868 --- /dev/null +++ b/model-00073-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:bf89244d0c55300a2767306b424eff449e2a515b505b917f9599006a8cd83dbb +size 8598786520 diff --git a/model-00074-of-000163.safetensors b/model-00074-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..4ca7f0aa10cb05e1e526793b8f049503ca857988 --- /dev/null +++ b/model-00074-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:6c560e37ac4f6be3c9824741a94cd63bb65b2f00655e9467734fdd3b99bc5e35 +size 8602554408 diff --git a/model-00075-of-000163.safetensors b/model-00075-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..8a35714a51ee1f821f74976c446f68e81406418f --- /dev/null +++ b/model-00075-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:692fff8d4f97a5ebe7bcb2dd5b4f7375fc4595b041b1480c15defa73fd95d8fa +size 8598786720 diff --git a/model-00076-of-000163.safetensors b/model-00076-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..4b8c4bb40300548ff046431632423ec1c8a6ff56 --- /dev/null +++ b/model-00076-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:033523de996f2982a34ab48c66d8cc2e49bfbec8fa7eabb1f71b3366d7845bfa +size 8602554208 diff --git a/model-00077-of-000163.safetensors b/model-00077-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..9f7177da0134bfb33cba7cbd9b6ba8bba9b9d703 --- /dev/null +++ b/model-00077-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:3400a940e5910e410faabf5448ffbb6d18d9f879834a9ba7d41510c4cd7b0f60 +size 8602554448 diff --git a/model-00078-of-000163.safetensors b/model-00078-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..4b4e1dc66f8515293fb6ee026a87e00ae36d66e8 --- /dev/null +++ b/model-00078-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:836e8b550e788541ec5ffb6b32ac4bae99978ddbd2d6f8fbcde5806f4b1b8e46 +size 3493899088 diff --git a/model-00080-of-000163.safetensors b/model-00080-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..5da161444f68d293bbc3b21da5cfc60f0c346acb --- /dev/null +++ b/model-00080-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:e21caf3b004f1d5bf52779d7b4226e97657223f2dd270a693325b7c43fcc2487 +size 8602554424 diff --git a/model-00081-of-000163.safetensors b/model-00081-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..4bc4cd9266834b26456a0dca7b720d9591fb4394 --- /dev/null +++ b/model-00081-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:268fdb4e0a94d29a884081ae75cf4472adf80dd50cda8583192bee09e9984d68 +size 8598786704 diff --git a/model-00082-of-000163.safetensors b/model-00082-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..7594e623651ea97730ba3c310672a8c5a069756b --- /dev/null +++ b/model-00082-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:1b3b165c4f92446bbebdccf65d60aad4ff0bba3e884c1e016e5f27b98fb7af54 +size 8602554224 diff --git a/model-00083-of-000163.safetensors b/model-00083-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..cf73eb2ba5aa2f8bedd426221e3f907c5a07a52a --- /dev/null +++ b/model-00083-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:09db84761d8e985cdd47d6366b798696c9ae54a478b21647b38845fea018272b +size 8602554448 diff --git a/model-00084-of-000163.safetensors b/model-00084-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..7621ae100a7edbe1886088ce41ab715e44b2f933 --- /dev/null +++ b/model-00084-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:46eb3c357ca882ec337412f7dbc6deef7914f2062ba72290d795eb5c8c22e765 +size 8598786608 diff --git a/model-00085-of-000163.safetensors b/model-00085-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..0ca64f7909452c3914b888a83d99dd54dc44f67c --- /dev/null +++ b/model-00085-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:a9e55ef86fe7535da527daeb5485fc689d7ca534afd87d484a2dd2bb00053e80 +size 8602554320 diff --git a/model-00086-of-000163.safetensors b/model-00086-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..61a47e6bd52e4836c0529761e00a5e6d88f900d8 --- /dev/null +++ b/model-00086-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:bb760c9de5041dd50b917d1b04fca66d5e61208b367d6798ad9d9d78f831dc5d +size 8602554448 diff --git a/model-00087-of-000163.safetensors b/model-00087-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..2106b8edd2f53fcccffab739bae1590373c2a4c5 --- /dev/null +++ b/model-00087-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:eff7f7a94c69b698c531f72dcef52e0c8e1da7d552b259d111977793f470a631 +size 8598786504 diff --git a/model-00088-of-000163.safetensors b/model-00088-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..e5f01661c275723c735aca8c9d1b136b706ce8c6 --- /dev/null +++ b/model-00088-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:32bfbd43255e0c0e89085b98515587ec832267dd78105c2bbfe5e158efb862f8 +size 8602554416 diff --git a/model-00089-of-000163.safetensors b/model-00089-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..2dd228eb609ed266f2ba71597288bdf69918753f --- /dev/null +++ b/model-00089-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:594549ec17c1c115fb616b8006b9dfa126993f6936860ec4cf290e36c130f9fa +size 8598786704 diff --git a/model-00090-of-000163.safetensors b/model-00090-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..3ee1df08273bb6257b7b7fa4ea9a047557081285 --- /dev/null +++ b/model-00090-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:af137cf1f8cca460b3642a839503e465883aff451fe7fd2fadb6e8a579e9185e +size 8602554224 diff --git a/model-00091-of-000163.safetensors b/model-00091-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..672318334e6298f7cbbb590b61faf945da026e39 --- /dev/null +++ b/model-00091-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:29492ca4644102ed8c38a512b23babfe5e58d6b08426d6e55e5efb2fa232d7e3 +size 8602554448 diff --git a/model-00092-of-000163.safetensors b/model-00092-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..aef9ef7d998a201597bb955c60d0be5c4f8ecadd --- /dev/null +++ b/model-00092-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:698d5cfb4b1043abb951fbc5450706e33a0d2da53baa19c0625da2f85470df35 +size 8598786616 diff --git a/model-00093-of-000163.safetensors b/model-00093-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..187177cae6e64e9c046f548dca2d432831c83e70 --- /dev/null +++ b/model-00093-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:e9049e18a0a8a65a03687e40749161cfe4b74874494ad474dc2dfd7457fd110d +size 8602554312 diff --git a/model-00094-of-000163.safetensors b/model-00094-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..1cb22cb8019e5f5946bf76200dd62d0d94506d85 --- /dev/null +++ b/model-00094-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:3b7dea3ac82289d7c805a708d5579d301c8d0068c1e05dbeb47535564fb9f4fb +size 8602554448 diff --git a/model-00095-of-000163.safetensors b/model-00095-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..ab0698b12f84a6965962b82ea1ca79d21cf9ef7b --- /dev/null +++ b/model-00095-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:f9273f1a3a4d1b24d5be1749f95f3f286832109362ed78df3903db865d694670 +size 8598786520 diff --git a/model-00096-of-000163.safetensors b/model-00096-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..83538e866e2043c5c00334cbacab1b7b99c9c1ee --- /dev/null +++ b/model-00096-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:1b939b1649e74279090c8e9c2b3a4baffa60e55372fa628b246ddf36f484de81 +size 8602554408 diff --git a/model-00098-of-000163.safetensors b/model-00098-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..9763363ae7bf3be401f365c7a2c26996e65a9671 --- /dev/null +++ b/model-00098-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:35102e570ed4de1f0660d5d6bfb8c367de2e2128e810f490b4ab84725f9e72c9 +size 8602554208 diff --git a/model-00099-of-000163.safetensors b/model-00099-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..6db9194ac961ebd4551e94e3233c2e16903e4894 --- /dev/null +++ b/model-00099-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:4c33ab4239693ecc54eb57bc188e1dcbd7dbe0db6e5bae1e61bfc1080bd392f9 +size 8602554448 diff --git a/model-00100-of-000163.safetensors b/model-00100-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..69dfaacc0ad284886ec8f299a128ce2bcdc67425 --- /dev/null +++ b/model-00100-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:cb6963e86fd900a84f4987feca7526361bb2e714e1ad1c428ebf9f1086722b78 +size 3493899088 diff --git a/model-00101-of-000163.safetensors b/model-00101-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..db1cd619158b173d8a9affbc9bd6afc2dfad10fc --- /dev/null +++ b/model-00101-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:928b263d92524520fa26dfd0b326b722eb43f333e513d61851207c039cdd413c +size 8598757608 diff --git a/model-00102-of-000163.safetensors b/model-00102-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..04964eb27bf5a22c94b37cbcafa4f61181de5c85 --- /dev/null +++ b/model-00102-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:730fcb83aa49dffd318543a55c4cdf8a1efacb3860e7ee6ade0344460ced2b59 +size 8602554424 diff --git a/model-00103-of-000163.safetensors b/model-00103-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..113262185540bea6406b7932f6f6e7978a275b49 --- /dev/null +++ b/model-00103-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:adc5c0d711e7e6c87f6dfc65872bfe897547ce21ed1bb96b8b5daad32becb2a3 +size 8598786704 diff --git a/model-00104-of-000163.safetensors b/model-00104-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..708ef618c8b48ee1d2a9b8b561187b9d37bc2404 --- /dev/null +++ b/model-00104-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:b653d66c811b7e851cb36afdfa39d77f353df1ec41e9a3d529bc8596ea64ddf4 +size 8602554224 diff --git a/model-00105-of-000163.safetensors b/model-00105-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..12abdb39d7f0853018858862a5293a58c1d1b6e0 --- /dev/null +++ b/model-00105-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:065f702fc487c7e31436f59168c740551d7bddfa0869d1a6fa43afa782c8d37e +size 8602554448 diff --git a/model-00106-of-000163.safetensors b/model-00106-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..d23a912bac46cb328ec420644c616df8a27efeb1 --- /dev/null +++ b/model-00106-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:a09deb76871c8369848a0b18ecbe8a064aa9e100ffa43dec0a36a2021aa69e73 +size 8598786608 diff --git a/model-00107-of-000163.safetensors b/model-00107-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..44fbe58340c3915f03547f2885f36fbe7196e13e --- /dev/null +++ b/model-00107-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:a758c0bd7c4930765ee695bf31d5a8ff43d8bf9ac46192f34e6372e37a302a44 +size 8602554320 diff --git a/model-00108-of-000163.safetensors b/model-00108-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..dd01c79d8c7677764447030c5410f529e05fda01 --- /dev/null +++ b/model-00108-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:9e9c2fac7b6df5b4d3d34bcb793741086b62c31f857b915ac32eab689a46433e +size 8602554448 diff --git a/model-00110-of-000163.safetensors b/model-00110-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..1483c8ffff6020071d3ebc0a6f539b52e2a39d26 --- /dev/null +++ b/model-00110-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:05cc592517d8b0b7dc436a6b8c0b754db466368b91ed392e8c3e608fbe721aeb +size 8602554416 diff --git a/model-00111-of-000163.safetensors b/model-00111-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..52fa612cc6cee8a7a1345e6114c61e49a85f91dd --- /dev/null +++ b/model-00111-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:693317186a4b52c0c5138da8ce6a647fae81e29aacc37c436122eea7cf0737f2 +size 8598786704 diff --git a/model-00112-of-000163.safetensors b/model-00112-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..d57f46bf7e61acf862813b759ea8ffa06f4a0be5 --- /dev/null +++ b/model-00112-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:166f049112924ff6e0aa8140918facd59a1e7fd02bda26fc9ae1d2e9a05dc8a8 +size 8602554224 diff --git a/model-00114-of-000163.safetensors b/model-00114-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..f03b575222503f88b8c810ef58af9575215e21a9 --- /dev/null +++ b/model-00114-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:3a4763e2e09fa783282106bb0d7f41c76913c0266ae761c058054ccfabcae2fe +size 8598786616 diff --git a/model-00115-of-000163.safetensors b/model-00115-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..c1885c8418703413e7391fa97223b4721cb41745 --- /dev/null +++ b/model-00115-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:d63e1cba331bffe5872449fbcc9402f8014b4ff7aa3f820e2b1e37a0111ebc74 +size 8602554312 diff --git a/model-00117-of-000163.safetensors b/model-00117-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..cee6bf98431a43ab9e5bcd485dcee69b908aae61 --- /dev/null +++ b/model-00117-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:71ebd04ec647608a85cd8d778aa643c27682e272f35b4328af1104456b759753 +size 8598786520 diff --git a/model-00118-of-000163.safetensors b/model-00118-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..2d4805e9bb1c728c3726a39e9fce4951f50dccff --- /dev/null +++ b/model-00118-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:0ac618b9c28cb85ffb68757bfa6e8b81b82fcae96ad6ad36133f15ee31ecca99 +size 8602554408 diff --git a/model-00119-of-000163.safetensors b/model-00119-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..94d7d2ba7429c5c10cb630a5d8bcd2c33fc7dd49 --- /dev/null +++ b/model-00119-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:216df7546806112b62ec8a4ecc3a8909c9da8089b35e7d05ba2d762e89c46231 +size 8598786720 diff --git a/model-00120-of-000163.safetensors b/model-00120-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..9b661215a5b586f30db7163968c49544220aeda1 --- /dev/null +++ b/model-00120-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:c9bd57f04c30f99834396efde60c26936e08c7c25aa74539cbec95d227c4d86d +size 8602554208 diff --git a/model-00121-of-000163.safetensors b/model-00121-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..11aee99ab4de7cc1f0a0d83625c68d77aa1c9dfd --- /dev/null +++ b/model-00121-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:77c0d256d8d8e237f6e212132560e4cf5d67f7e24f4fab14a0b0bbc687cf7172 +size 8602554448 diff --git a/model-00122-of-000163.safetensors b/model-00122-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..0743df33ee47a850bded8a636bd5afa5dbcff127 --- /dev/null +++ b/model-00122-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:85e9e443d9adc3c25eeed1d136646a77d6f1b9abac5142338d34928514e07bb4 +size 3493899088 diff --git a/model-00123-of-000163.safetensors b/model-00123-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..d3e567bdd07368657c8f4ff4ecf71e2504d0c848 --- /dev/null +++ b/model-00123-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:00123b9d32302b36ce51804d43d21d4b5d92e30f4746b04c0d2f6f71f0d0e3f0 +size 8598757608 diff --git a/model-00124-of-000163.safetensors b/model-00124-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..7129876b3938211eae41616b039e8bdcfe2a93ec --- /dev/null +++ b/model-00124-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:4d98f426ac234d180c2e8cbc7c3ccf8a2762e78bf1e4d39474cc5fe147b56d5e +size 8602554424 diff --git a/model-00125-of-000163.safetensors b/model-00125-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..cbf6b5a85294639a60a2027a4dec14c33c2765ff --- /dev/null +++ b/model-00125-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:ad5aa9246b3c5cd2c576527e1a8358af565b58663738ff5ac439f0679b3b2f15 +size 8598786704 diff --git a/model-00126-of-000163.safetensors b/model-00126-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..65edb62d14651aeb136d8c483d92d80aee2bf396 --- /dev/null +++ b/model-00126-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:880b53640cf8b75da9b828b93c373e63e6b613f491e940a15afd39d25f030fc0 +size 8602554224 diff --git a/model-00127-of-000163.safetensors b/model-00127-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..e8274672d35ef65ef6ec5a50717b10ea30ac0d8e --- /dev/null +++ b/model-00127-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:b64ae7cbdccb7e284b92cdcc2929e8201213158444c91063d370a075012c59a8 +size 8602554448 diff --git a/model-00128-of-000163.safetensors b/model-00128-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..9dc47ac76a1b891267c3134413631abc0bdbac82 --- /dev/null +++ b/model-00128-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:e2c05fa9b521b04f64d2776b01ea39f6f0329571212b684e8565b446a55ddc69 +size 8598786608 diff --git a/model-00129-of-000163.safetensors b/model-00129-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..b9b2461465c9531874eb1ac7f464c997bb417636 --- /dev/null +++ b/model-00129-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:acef59a42ab90ba525a9e0c4a985c09aa8505a35c43fe05a1d27055dc092d6f9 +size 8602554320 diff --git a/model-00130-of-000163.safetensors b/model-00130-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..a80d7bb39c4f5271fd801ec3f083114d76555fb4 --- /dev/null +++ b/model-00130-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:191efcb688260af5878ec6dc2882305924095cbd623e5eaf6f29f31013dd098c +size 8602554448 diff --git a/model-00131-of-000163.safetensors b/model-00131-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..1114bcd10670ead8e589bfd3c2ed2fa9bef56dac --- /dev/null +++ b/model-00131-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:0e2e44df80ba461143aa992e33b1b07662650c5a81ce3e9f73c4045272c906d4 +size 8598786504 diff --git a/model-00132-of-000163.safetensors b/model-00132-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..ebe5a23cb53bbdfbbc1c4814d3c377f9e06a3fac --- /dev/null +++ b/model-00132-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:4f174137e0c267b316ca3e13655a37f583bf90605d85083dfe320ca26fb02d7c +size 8602554416 diff --git a/model-00133-of-000163.safetensors b/model-00133-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..8db45f04e194a1428a10a452b2f3bf984d859d8c --- /dev/null +++ b/model-00133-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:422a29d69bb14b76541fbfca661ced0016130485dabba0c2b3f4198411a40dc1 +size 8598786704 diff --git a/model-00134-of-000163.safetensors b/model-00134-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..24a49a344e549bf98753b5e95e257f7db6c8f5f2 --- /dev/null +++ b/model-00134-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:e888344a564ac24700d6b4fb1d7e42aabdffd309847cafc908bf74e7b39e2fda +size 8602554224 diff --git a/model-00135-of-000163.safetensors b/model-00135-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..896e28047b4cfa79091a89e7bb574f025f7209ef --- /dev/null +++ b/model-00135-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:45e465ad59f63fdb68eaf4ca2850c881a6d2f776e1c2c668c4b771acd2a8fed9 +size 8602554448 diff --git a/model-00136-of-000163.safetensors b/model-00136-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..bda5e541f984800a6c0c45e6720745b10889969c --- /dev/null +++ b/model-00136-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:af5c701436337be68fa83f9616669513df08dda6e5aa047962cb41b82003f2fc +size 8598786616 diff --git a/model-00137-of-000163.safetensors b/model-00137-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..383c2853e604f7f9c23784911c3c5538edc8cf5e --- /dev/null +++ b/model-00137-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:a5d8c952eec65fd8dd86e14a57a6e356c40c8f82b5e156627a82327a4b6b6031 +size 8602554312 diff --git a/model-00139-of-000163.safetensors b/model-00139-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..ded5649b0595df3d5e4b48b5e95503e3023b917f --- /dev/null +++ b/model-00139-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:08dc5eacfd88c22b84cafc5578b26882a4cc0e5ea9937919ffc9c795a3676a4c +size 8598786520 diff --git a/model-00140-of-000163.safetensors b/model-00140-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..e940f4dea60aaca040fbc10422b7bc08d4c2b3ff --- /dev/null +++ b/model-00140-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:e481555e5e388d5123a3462ac9204af71def389fbc02f83eac0c84bd0404fc67 +size 8602554408 diff --git a/model-00141-of-000163.safetensors b/model-00141-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..c4b97d336893fbe701cb50ea836008f95e00549e --- /dev/null +++ b/model-00141-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:fd97d1767f887dc24aeeae4aba07b111d64cd4294ea7ad9940764b0f1cca7971 +size 6283123256 diff --git a/model-00142-of-000163.safetensors b/model-00142-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..57b1096c3c6a3a1c275faa3477a88c9fc43c919e --- /dev/null +++ b/model-00142-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:abfa094e686bd0f92345dacfaf1125eec7b17f901f690a18e680da93f8cbad74 +size 8598757608 diff --git a/model-00143-of-000163.safetensors b/model-00143-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..182b8550999a68d550a48c63a5ce0d912ca04ebe --- /dev/null +++ b/model-00143-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:6e933e95e9561fa3fb7e6f44d6af3ea36bb268f7ddf0a3201f48ab09183cfccf +size 8602554424 diff --git a/model-00145-of-000163.safetensors b/model-00145-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..f0afaaad5ea98e58de3e046c1cf7f70c5ae50578 --- /dev/null +++ b/model-00145-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:98e771c1520435524f1f2b13af42f8d425aaa646c54eeaf54bc370d15290c203 +size 8602554224 diff --git a/model-00146-of-000163.safetensors b/model-00146-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..b3758b47ecf7244d80034ba9e1fdb39117e4c1ad --- /dev/null +++ b/model-00146-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:cd26d981cf45884d305afca916e1ee10c28978113436f4b183d9928d9ab3ccca +size 8602554448 diff --git a/model-00147-of-000163.safetensors b/model-00147-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..ae40fdcf184d211283b23188b7e3369808dbd7de --- /dev/null +++ b/model-00147-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:82d4da6025951d81f936fd7c99afa65e4daf302fce894794ff799a355aa55672 +size 8598786608 diff --git a/model-00148-of-000163.safetensors b/model-00148-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..e74341f71d070d2864953a1e6aa57549be792822 --- /dev/null +++ b/model-00148-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:22bc8c994b5574072e40b30937d38d7028726abde67e24b5b2f13901d65110c6 +size 8602554320 diff --git a/model-00149-of-000163.safetensors b/model-00149-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..572724ca1b3c179eca22f62c487ccdc9ed255646 --- /dev/null +++ b/model-00149-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:65facc6a621d82a2af92e03d0635dfc7d1affc501fece14af394be58b8fb9de8 +size 8602554448 diff --git a/model-00150-of-000163.safetensors b/model-00150-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..8be0f239c0c25fc65df777450b75fac3256e5f81 --- /dev/null +++ b/model-00150-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:5d3df7be3bf159f6c31a86e34b2115d0044e3391301f1fbc5e641285f2d5293e +size 8598786504 diff --git a/model-00152-of-000163.safetensors b/model-00152-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..025bb79872e6978e0ea9ecd47df28de51600155f --- /dev/null +++ b/model-00152-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:70587dcd6fdae5681fcba0d0aa19ddaf175d3feef2d5c3e47b8f8e6fab0e7167 +size 8598786704 diff --git a/model-00153-of-000163.safetensors b/model-00153-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..a47bf37b226a61f0bc5f0e04309afee40e0ff8f5 --- /dev/null +++ b/model-00153-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:e13964da608a929d25f7d0a6b5b850d1221d446f6bf58bea00a23ceaf77bd192 +size 8602554224 diff --git a/model-00154-of-000163.safetensors b/model-00154-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..0e0b92e20659acefb2b97cb0c986e74ade132a09 --- /dev/null +++ b/model-00154-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:7aeb9c130a2ff2d15bce05cd3318b095d2ad604258d3b5bae28437d594dc8dc1 +size 8602554448 diff --git a/model-00156-of-000163.safetensors b/model-00156-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..641a34561df699423abb85993a27b04219c22f5a --- /dev/null +++ b/model-00156-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:31393ee79e481493b72b3f23f5ca4fd90d1180cab0b105fab4de681106cca888 +size 8602554312 diff --git a/model-00157-of-000163.safetensors b/model-00157-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..e246e74b071c3996700aea1cf9b63ff8739b56d5 --- /dev/null +++ b/model-00157-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:af61ad2dc4c87cd59d4b3668cf8d68979512d4ef5aab3104da7b6118a14a9c00 +size 8602554448 diff --git a/model-00158-of-000163.safetensors b/model-00158-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..fc7d4668c8d64736496414100cd825d52117b7a0 --- /dev/null +++ b/model-00158-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:2d2f56a2513e7cb200fed07531addfb881a488bd5e3b2e661d7c28830277b056 +size 8598786520 diff --git a/model-00160-of-000163.safetensors b/model-00160-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..4192f237548ec1884ba5c30f39393192f3d9a541 --- /dev/null +++ b/model-00160-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:6c8a932ce49bba0f5489f30a13708532cb8bc727667dd76ac60a4371377095ef +size 8602463472 diff --git a/model-00161-of-000163.safetensors b/model-00161-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..b23307755a9457a4a9914a78854a60831ae86ee4 --- /dev/null +++ b/model-00161-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:b65a6d2704cdc7c20c5e7824e8c0461e68adecc97977b437a00d5d77b8424d24 +size 8602554128 diff --git a/model-00162-of-000163.safetensors b/model-00162-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..969df4465d0ae93b80862cade1c9ddf9a744469c --- /dev/null +++ b/model-00162-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:0f8c12d03556374bdbc47ada0c4e7d299a4615dbd2d74da07218ac55b35a76e6 +size 8602554440 diff --git a/model-00163-of-000163.safetensors b/model-00163-of-000163.safetensors new file mode 100644 index 0000000000000000000000000000000000000000..123294428c1603b579b958ce202aebcfed7adbd5 --- /dev/null +++ b/model-00163-of-000163.safetensors @@ -0,0 +1,3 @@ +version https://git-lfs.github.com/spec/v1 +oid sha256:8bc6106794d6f881eae142a1de5ef073ce0beb3204a4722de547dd72d606bd42 +size 9255875920 diff --git a/model.safetensors.index.json b/model.safetensors.index.json new file mode 100644 index 0000000000000000000000000000000000000000..5b60276c1efb766f6a1aee2e357e27407f6ed12b --- /dev/null +++ b/model.safetensors.index.json @@ -0,0 +1,46188 @@ +{ + "metadata": {}, + "weight_map": { + "model.embed_tokens.weight": "model-00001-of-000163.safetensors", + "model.layers.0.self_attn.q_a_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.0.self_attn.q_a_layernorm.weight": "model-00001-of-000163.safetensors", + "model.layers.0.self_attn.q_b_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.0.self_attn.kv_a_proj_with_mqa.weight": "model-00001-of-000163.safetensors", + "model.layers.0.self_attn.kv_a_layernorm.weight": "model-00001-of-000163.safetensors", + "model.layers.0.self_attn.kv_b_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.0.self_attn.o_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.0.mlp.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.0.mlp.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.0.mlp.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.0.input_layernorm.weight": "model-00001-of-000163.safetensors", + "model.layers.0.post_attention_layernorm.weight": "model-00001-of-000163.safetensors", + "model.layers.1.self_attn.q_a_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.1.self_attn.q_a_layernorm.weight": "model-00001-of-000163.safetensors", + "model.layers.1.self_attn.q_b_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.1.self_attn.kv_a_proj_with_mqa.weight": "model-00001-of-000163.safetensors", + "model.layers.1.self_attn.kv_a_layernorm.weight": "model-00001-of-000163.safetensors", + "model.layers.1.self_attn.kv_b_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.1.self_attn.o_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.1.mlp.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.1.mlp.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.1.mlp.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.1.input_layernorm.weight": "model-00001-of-000163.safetensors", + "model.layers.1.post_attention_layernorm.weight": "model-00001-of-000163.safetensors", + "model.layers.2.self_attn.q_a_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.2.self_attn.q_a_layernorm.weight": "model-00001-of-000163.safetensors", + "model.layers.2.self_attn.q_b_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.2.self_attn.kv_a_proj_with_mqa.weight": "model-00001-of-000163.safetensors", + "model.layers.2.self_attn.kv_a_layernorm.weight": "model-00001-of-000163.safetensors", + "model.layers.2.self_attn.kv_b_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.2.self_attn.o_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.2.mlp.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.2.mlp.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.2.mlp.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.2.input_layernorm.weight": "model-00001-of-000163.safetensors", + "model.layers.2.post_attention_layernorm.weight": "model-00001-of-000163.safetensors", + "model.layers.3.self_attn.q_a_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.self_attn.q_a_layernorm.weight": "model-00001-of-000163.safetensors", + "model.layers.3.self_attn.q_b_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.self_attn.kv_a_proj_with_mqa.weight": "model-00001-of-000163.safetensors", + "model.layers.3.self_attn.kv_a_layernorm.weight": "model-00001-of-000163.safetensors", + "model.layers.3.self_attn.kv_b_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.self_attn.o_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.gate.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.gate.e_score_correction_bias": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.shared_experts.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.shared_experts.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.shared_experts.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.0.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.0.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.0.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.1.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.1.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.1.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.2.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.2.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.2.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.3.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.3.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.3.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.4.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.4.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.4.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.5.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.5.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.5.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.6.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.6.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.6.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.7.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.7.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.7.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.8.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.8.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.8.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.9.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.9.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.9.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.10.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.10.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.10.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.11.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.11.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.11.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.12.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.12.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.12.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.13.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.13.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.13.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.14.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.14.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.14.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.15.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.15.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.15.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.16.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.16.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.16.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.17.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.17.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.17.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.18.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.18.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.18.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.19.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.19.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.19.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.20.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.20.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.20.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.21.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.21.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.21.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.22.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.22.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.22.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.23.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.23.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.23.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.24.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.24.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.24.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.25.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.25.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.25.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.26.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.26.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.26.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.27.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.27.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.27.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.28.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.28.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.28.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.29.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.29.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.29.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.30.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.30.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.30.down_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.31.gate_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.31.up_proj.weight": "model-00001-of-000163.safetensors", + "model.layers.3.mlp.experts.31.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.32.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.32.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.32.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.33.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.33.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.33.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.34.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.34.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.34.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.35.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.35.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.35.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.36.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.36.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.36.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.37.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.37.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.37.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.38.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.38.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.38.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.39.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.39.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.39.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.40.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.40.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.40.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.41.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.41.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.41.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.42.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.42.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.42.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.43.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.43.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.43.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.44.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.44.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.44.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.45.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.45.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.45.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.46.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.46.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.46.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.47.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.47.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.47.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.48.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.48.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.48.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.49.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.49.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.49.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.50.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.50.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.50.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.51.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.51.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.51.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.52.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.52.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.52.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.53.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.53.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.53.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.54.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.54.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.54.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.55.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.55.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.55.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.56.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.56.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.56.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.57.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.57.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.57.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.58.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.58.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.58.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.59.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.59.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.59.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.60.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.60.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.60.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.61.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.61.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.61.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.62.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.62.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.62.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.63.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.63.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.63.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.64.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.64.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.64.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.65.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.65.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.65.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.66.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.66.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.66.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.67.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.67.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.67.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.68.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.68.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.68.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.69.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.69.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.69.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.70.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.70.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.70.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.71.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.71.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.71.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.72.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.72.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.72.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.73.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.73.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.73.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.74.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.74.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.74.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.75.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.75.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.75.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.76.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.76.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.76.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.77.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.77.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.77.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.78.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.78.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.78.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.79.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.79.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.79.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.80.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.80.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.80.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.81.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.81.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.81.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.82.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.82.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.82.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.83.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.83.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.83.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.84.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.84.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.84.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.85.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.85.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.85.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.86.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.86.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.86.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.87.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.87.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.87.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.88.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.88.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.88.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.89.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.89.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.89.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.90.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.90.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.90.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.91.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.91.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.91.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.92.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.92.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.92.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.93.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.93.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.93.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.94.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.94.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.94.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.95.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.95.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.95.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.96.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.96.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.96.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.97.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.97.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.97.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.98.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.98.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.98.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.99.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.99.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.99.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.100.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.100.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.100.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.101.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.101.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.101.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.102.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.102.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.102.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.103.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.103.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.103.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.104.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.104.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.104.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.105.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.105.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.105.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.106.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.106.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.106.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.107.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.107.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.107.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.108.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.108.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.108.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.109.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.109.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.109.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.110.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.110.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.110.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.111.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.111.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.111.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.112.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.112.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.112.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.113.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.113.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.113.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.114.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.114.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.114.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.115.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.115.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.115.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.116.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.116.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.116.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.117.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.117.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.117.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.118.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.118.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.118.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.119.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.119.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.119.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.120.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.120.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.120.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.121.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.121.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.121.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.122.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.122.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.122.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.123.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.123.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.123.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.124.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.124.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.124.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.125.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.125.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.125.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.126.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.126.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.126.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.127.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.127.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.127.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.128.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.128.up_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.128.down_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.129.gate_proj.weight": "model-00002-of-000163.safetensors", + "model.layers.3.mlp.experts.129.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.129.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.130.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.130.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.130.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.131.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.131.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.131.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.132.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.132.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.132.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.133.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.133.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.133.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.134.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.134.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.134.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.135.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.135.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.135.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.136.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.136.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.136.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.137.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.137.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.137.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.138.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.138.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.138.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.139.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.139.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.139.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.140.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.140.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.140.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.141.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.141.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.141.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.142.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.142.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.142.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.143.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.143.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.143.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.144.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.144.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.144.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.145.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.145.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.145.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.146.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.146.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.146.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.147.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.147.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.147.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.148.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.148.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.148.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.149.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.149.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.149.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.150.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.150.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.150.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.151.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.151.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.151.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.152.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.152.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.152.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.153.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.153.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.153.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.154.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.154.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.154.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.155.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.155.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.155.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.156.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.156.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.156.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.157.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.157.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.157.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.158.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.158.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.158.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.159.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.159.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.159.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.160.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.160.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.160.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.161.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.161.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.161.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.162.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.162.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.162.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.163.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.163.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.163.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.164.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.164.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.164.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.165.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.165.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.165.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.166.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.166.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.166.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.167.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.167.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.167.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.168.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.168.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.168.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.169.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.169.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.169.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.170.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.170.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.170.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.171.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.171.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.171.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.172.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.172.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.172.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.173.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.173.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.173.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.174.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.174.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.174.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.175.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.175.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.175.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.176.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.176.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.176.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.177.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.177.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.177.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.178.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.178.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.178.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.179.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.179.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.179.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.180.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.180.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.180.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.181.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.181.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.181.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.182.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.182.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.182.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.183.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.183.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.183.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.184.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.184.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.184.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.185.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.185.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.185.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.186.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.186.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.186.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.187.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.187.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.187.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.188.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.188.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.188.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.189.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.189.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.189.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.190.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.190.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.190.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.191.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.191.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.191.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.192.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.192.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.192.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.193.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.193.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.193.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.194.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.194.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.194.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.195.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.195.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.195.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.196.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.196.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.196.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.197.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.197.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.197.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.198.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.198.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.198.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.199.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.199.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.199.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.200.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.200.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.200.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.201.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.201.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.201.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.202.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.202.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.202.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.203.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.203.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.203.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.204.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.204.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.204.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.205.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.205.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.205.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.206.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.206.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.206.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.207.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.207.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.207.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.208.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.208.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.208.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.209.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.209.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.209.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.210.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.210.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.210.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.211.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.211.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.211.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.212.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.212.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.212.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.213.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.213.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.213.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.214.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.214.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.214.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.215.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.215.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.215.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.216.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.216.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.216.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.217.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.217.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.217.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.218.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.218.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.218.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.219.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.219.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.219.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.220.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.220.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.220.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.221.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.221.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.221.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.222.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.222.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.222.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.223.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.223.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.223.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.224.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.224.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.224.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.225.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.225.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.225.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.226.gate_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.226.up_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.226.down_proj.weight": "model-00003-of-000163.safetensors", + "model.layers.3.mlp.experts.227.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.227.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.227.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.228.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.228.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.228.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.229.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.229.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.229.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.230.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.230.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.230.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.231.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.231.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.231.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.232.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.232.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.232.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.233.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.233.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.233.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.234.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.234.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.234.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.235.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.235.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.235.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.236.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.236.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.236.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.237.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.237.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.237.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.238.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.238.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.238.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.239.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.239.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.239.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.240.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.240.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.240.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.241.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.241.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.241.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.242.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.242.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.242.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.243.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.243.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.243.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.244.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.244.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.244.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.245.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.245.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.245.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.246.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.246.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.246.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.247.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.247.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.247.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.248.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.248.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.248.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.249.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.249.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.249.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.250.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.250.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.250.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.251.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.251.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.251.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.252.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.252.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.252.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.253.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.253.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.253.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.254.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.254.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.254.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.255.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.255.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.mlp.experts.255.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.3.input_layernorm.weight": "model-00004-of-000163.safetensors", + "model.layers.3.post_attention_layernorm.weight": "model-00004-of-000163.safetensors", + "model.layers.4.self_attn.q_a_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.self_attn.q_a_layernorm.weight": "model-00004-of-000163.safetensors", + "model.layers.4.self_attn.q_b_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.self_attn.kv_a_proj_with_mqa.weight": "model-00004-of-000163.safetensors", + "model.layers.4.self_attn.kv_a_layernorm.weight": "model-00004-of-000163.safetensors", + "model.layers.4.self_attn.kv_b_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.self_attn.o_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.gate.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.gate.e_score_correction_bias": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.shared_experts.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.shared_experts.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.shared_experts.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.0.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.0.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.0.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.1.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.1.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.1.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.2.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.2.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.2.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.3.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.3.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.3.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.4.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.4.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.4.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.5.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.5.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.5.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.6.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.6.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.6.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.7.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.7.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.7.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.8.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.8.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.8.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.9.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.9.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.9.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.10.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.10.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.10.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.11.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.11.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.11.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.12.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.12.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.12.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.13.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.13.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.13.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.14.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.14.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.14.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.15.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.15.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.15.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.16.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.16.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.16.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.17.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.17.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.17.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.18.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.18.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.18.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.19.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.19.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.19.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.20.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.20.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.20.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.21.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.21.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.21.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.22.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.22.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.22.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.23.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.23.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.23.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.24.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.24.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.24.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.25.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.25.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.25.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.26.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.26.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.26.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.27.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.27.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.27.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.28.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.28.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.28.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.29.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.29.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.29.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.30.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.30.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.30.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.31.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.31.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.31.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.32.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.32.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.32.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.33.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.33.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.33.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.34.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.34.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.34.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.35.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.35.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.35.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.36.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.36.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.36.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.37.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.37.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.37.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.38.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.38.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.38.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.39.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.39.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.39.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.40.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.40.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.40.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.41.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.41.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.41.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.42.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.42.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.42.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.43.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.43.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.43.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.44.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.44.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.44.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.45.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.45.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.45.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.46.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.46.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.46.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.47.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.47.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.47.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.48.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.48.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.48.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.49.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.49.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.49.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.50.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.50.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.50.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.51.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.51.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.51.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.52.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.52.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.52.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.53.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.53.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.53.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.54.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.54.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.54.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.55.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.55.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.55.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.56.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.56.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.56.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.57.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.57.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.57.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.58.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.58.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.58.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.59.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.59.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.59.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.60.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.60.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.60.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.61.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.61.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.61.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.62.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.62.up_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.62.down_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.63.gate_proj.weight": "model-00004-of-000163.safetensors", + "model.layers.4.mlp.experts.63.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.63.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.64.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.64.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.64.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.65.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.65.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.65.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.66.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.66.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.66.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.67.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.67.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.67.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.68.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.68.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.68.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.69.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.69.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.69.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.70.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.70.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.70.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.71.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.71.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.71.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.72.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.72.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.72.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.73.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.73.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.73.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.74.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.74.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.74.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.75.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.75.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.75.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.76.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.76.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.76.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.77.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.77.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.77.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.78.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.78.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.78.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.79.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.79.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.79.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.80.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.80.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.80.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.81.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.81.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.81.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.82.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.82.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.82.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.83.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.83.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.83.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.84.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.84.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.84.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.85.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.85.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.85.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.86.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.86.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.86.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.87.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.87.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.87.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.88.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.88.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.88.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.89.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.89.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.89.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.90.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.90.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.90.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.91.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.91.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.91.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.92.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.92.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.92.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.93.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.93.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.93.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.94.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.94.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.94.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.95.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.95.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.95.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.96.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.96.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.96.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.97.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.97.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.97.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.98.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.98.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.98.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.99.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.99.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.99.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.100.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.100.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.100.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.101.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.101.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.101.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.102.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.102.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.102.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.103.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.103.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.103.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.104.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.104.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.104.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.105.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.105.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.105.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.106.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.106.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.106.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.107.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.107.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.107.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.108.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.108.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.108.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.109.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.109.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.109.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.110.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.110.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.110.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.111.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.111.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.111.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.112.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.112.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.112.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.113.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.113.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.113.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.114.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.114.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.114.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.115.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.115.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.115.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.116.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.116.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.116.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.117.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.117.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.117.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.118.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.118.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.118.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.119.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.119.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.119.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.120.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.120.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.120.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.121.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.121.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.121.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.122.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.122.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.122.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.123.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.123.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.123.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.124.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.124.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.124.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.125.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.125.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.125.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.126.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.126.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.126.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.127.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.127.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.127.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.128.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.128.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.128.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.129.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.129.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.129.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.130.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.130.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.130.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.131.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.131.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.131.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.132.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.132.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.132.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.133.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.133.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.133.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.134.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.134.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.134.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.135.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.135.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.135.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.136.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.136.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.136.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.137.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.137.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.137.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.138.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.138.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.138.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.139.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.139.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.139.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.140.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.140.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.140.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.141.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.141.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.141.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.142.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.142.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.142.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.143.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.143.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.143.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.144.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.144.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.144.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.145.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.145.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.145.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.146.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.146.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.146.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.147.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.147.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.147.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.148.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.148.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.148.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.149.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.149.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.149.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.150.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.150.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.150.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.151.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.151.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.151.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.152.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.152.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.152.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.153.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.153.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.153.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.154.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.154.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.154.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.155.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.155.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.155.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.156.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.156.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.156.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.157.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.157.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.157.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.158.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.158.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.158.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.159.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.159.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.159.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.160.gate_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.160.up_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.160.down_proj.weight": "model-00005-of-000163.safetensors", + "model.layers.4.mlp.experts.161.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.161.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.161.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.162.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.162.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.162.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.163.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.163.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.163.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.164.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.164.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.164.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.165.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.165.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.165.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.166.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.166.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.166.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.167.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.167.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.167.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.168.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.168.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.168.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.169.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.169.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.169.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.170.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.170.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.170.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.171.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.171.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.171.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.172.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.172.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.172.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.173.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.173.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.173.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.174.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.174.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.174.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.175.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.175.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.175.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.176.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.176.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.176.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.177.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.177.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.177.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.178.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.178.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.178.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.179.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.179.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.179.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.180.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.180.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.180.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.181.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.181.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.181.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.182.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.182.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.182.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.183.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.183.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.183.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.184.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.184.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.184.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.185.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.185.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.185.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.186.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.186.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.186.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.187.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.187.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.187.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.188.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.188.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.188.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.189.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.189.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.189.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.190.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.190.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.190.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.191.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.191.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.191.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.192.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.192.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.192.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.193.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.193.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.193.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.194.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.194.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.194.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.195.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.195.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.195.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.196.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.196.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.196.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.197.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.197.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.197.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.198.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.198.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.198.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.199.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.199.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.199.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.200.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.200.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.200.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.201.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.201.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.201.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.202.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.202.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.202.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.203.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.203.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.203.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.204.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.204.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.204.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.205.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.205.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.205.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.206.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.206.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.206.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.207.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.207.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.207.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.208.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.208.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.208.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.209.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.209.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.209.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.210.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.210.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.210.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.211.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.211.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.211.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.212.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.212.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.212.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.213.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.213.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.213.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.214.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.214.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.214.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.215.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.215.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.215.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.216.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.216.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.216.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.217.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.217.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.217.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.218.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.218.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.218.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.219.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.219.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.219.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.220.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.220.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.220.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.221.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.221.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.221.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.222.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.222.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.222.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.223.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.223.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.223.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.224.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.224.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.224.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.225.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.225.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.225.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.226.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.226.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.226.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.227.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.227.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.227.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.228.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.228.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.228.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.229.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.229.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.229.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.230.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.230.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.230.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.231.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.231.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.231.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.232.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.232.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.232.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.233.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.233.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.233.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.234.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.234.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.234.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.235.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.235.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.235.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.236.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.236.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.236.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.237.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.237.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.237.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.238.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.238.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.238.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.239.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.239.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.239.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.240.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.240.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.240.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.241.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.241.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.241.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.242.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.242.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.242.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.243.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.243.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.243.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.244.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.244.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.244.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.245.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.245.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.245.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.246.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.246.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.246.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.247.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.247.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.247.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.248.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.248.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.248.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.249.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.249.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.249.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.250.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.250.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.250.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.251.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.251.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.251.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.252.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.252.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.252.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.253.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.253.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.253.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.254.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.254.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.254.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.255.gate_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.255.up_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.mlp.experts.255.down_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.4.input_layernorm.weight": "model-00006-of-000163.safetensors", + "model.layers.4.post_attention_layernorm.weight": "model-00006-of-000163.safetensors", + "model.layers.5.self_attn.q_a_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.5.self_attn.q_a_layernorm.weight": "model-00006-of-000163.safetensors", + "model.layers.5.self_attn.q_b_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.5.self_attn.kv_a_proj_with_mqa.weight": "model-00006-of-000163.safetensors", + "model.layers.5.self_attn.kv_a_layernorm.weight": "model-00006-of-000163.safetensors", + "model.layers.5.self_attn.kv_b_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.5.self_attn.o_proj.weight": "model-00006-of-000163.safetensors", + "model.layers.5.mlp.gate.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.gate.e_score_correction_bias": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.shared_experts.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.shared_experts.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.shared_experts.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.0.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.0.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.0.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.1.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.1.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.1.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.2.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.2.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.2.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.3.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.3.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.3.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.4.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.4.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.4.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.5.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.5.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.5.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.6.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.6.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.6.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.7.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.7.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.7.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.8.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.8.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.8.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.9.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.9.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.9.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.10.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.10.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.10.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.11.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.11.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.11.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.12.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.12.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.12.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.13.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.13.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.13.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.14.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.14.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.14.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.15.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.15.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.15.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.16.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.16.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.16.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.17.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.17.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.17.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.18.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.18.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.18.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.19.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.19.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.19.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.20.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.20.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.20.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.21.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.21.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.21.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.22.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.22.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.22.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.23.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.23.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.23.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.24.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.24.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.24.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.25.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.25.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.25.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.26.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.26.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.26.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.27.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.27.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.27.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.28.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.28.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.28.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.29.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.29.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.29.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.30.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.30.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.30.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.31.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.31.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.31.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.32.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.32.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.32.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.33.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.33.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.33.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.34.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.34.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.34.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.35.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.35.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.35.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.36.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.36.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.36.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.37.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.37.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.37.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.38.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.38.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.38.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.39.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.39.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.39.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.40.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.40.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.40.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.41.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.41.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.41.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.42.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.42.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.42.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.43.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.43.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.43.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.44.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.44.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.44.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.45.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.45.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.45.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.46.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.46.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.46.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.47.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.47.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.47.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.48.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.48.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.48.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.49.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.49.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.49.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.50.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.50.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.50.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.51.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.51.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.51.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.52.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.52.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.52.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.53.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.53.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.53.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.54.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.54.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.54.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.55.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.55.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.55.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.56.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.56.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.56.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.57.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.57.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.57.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.58.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.58.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.58.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.59.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.59.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.59.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.60.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.60.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.60.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.61.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.61.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.61.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.62.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.62.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.62.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.63.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.63.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.63.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.64.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.64.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.64.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.65.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.65.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.65.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.66.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.66.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.66.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.67.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.67.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.67.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.68.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.68.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.68.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.69.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.69.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.69.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.70.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.70.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.70.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.71.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.71.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.71.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.72.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.72.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.72.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.73.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.73.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.73.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.74.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.74.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.74.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.75.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.75.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.75.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.76.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.76.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.76.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.77.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.77.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.77.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.78.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.78.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.78.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.79.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.79.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.79.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.80.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.80.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.80.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.81.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.81.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.81.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.82.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.82.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.82.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.83.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.83.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.83.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.84.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.84.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.84.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.85.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.85.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.85.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.86.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.86.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.86.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.87.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.87.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.87.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.88.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.88.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.88.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.89.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.89.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.89.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.90.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.90.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.90.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.91.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.91.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.91.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.92.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.92.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.92.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.93.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.93.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.93.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.94.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.94.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.94.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.95.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.95.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.95.down_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.96.gate_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.96.up_proj.weight": "model-00007-of-000163.safetensors", + "model.layers.5.mlp.experts.96.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.97.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.97.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.97.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.98.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.98.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.98.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.99.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.99.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.99.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.100.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.100.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.100.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.101.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.101.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.101.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.102.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.102.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.102.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.103.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.103.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.103.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.104.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.104.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.104.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.105.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.105.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.105.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.106.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.106.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.106.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.107.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.107.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.107.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.108.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.108.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.108.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.109.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.109.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.109.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.110.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.110.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.110.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.111.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.111.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.111.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.112.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.112.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.112.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.113.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.113.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.113.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.114.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.114.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.114.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.115.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.115.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.115.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.116.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.116.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.116.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.117.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.117.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.117.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.118.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.118.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.118.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.119.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.119.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.119.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.120.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.120.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.120.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.121.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.121.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.121.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.122.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.122.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.122.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.123.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.123.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.123.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.124.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.124.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.124.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.125.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.125.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.125.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.126.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.126.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.126.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.127.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.127.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.127.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.128.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.128.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.128.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.129.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.129.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.129.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.130.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.130.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.130.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.131.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.131.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.131.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.132.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.132.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.132.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.133.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.133.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.133.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.134.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.134.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.134.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.135.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.135.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.135.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.136.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.136.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.136.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.137.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.137.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.137.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.138.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.138.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.138.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.139.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.139.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.139.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.140.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.140.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.140.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.141.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.141.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.141.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.142.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.142.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.142.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.143.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.143.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.143.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.144.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.144.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.144.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.145.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.145.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.145.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.146.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.146.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.146.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.147.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.147.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.147.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.148.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.148.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.148.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.149.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.149.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.149.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.150.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.150.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.150.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.151.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.151.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.151.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.152.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.152.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.152.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.153.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.153.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.153.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.154.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.154.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.154.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.155.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.155.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.155.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.156.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.156.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.156.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.157.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.157.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.157.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.158.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.158.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.158.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.159.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.159.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.159.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.160.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.160.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.160.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.161.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.161.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.161.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.162.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.162.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.162.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.163.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.163.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.163.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.164.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.164.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.164.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.165.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.165.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.165.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.166.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.166.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.166.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.167.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.167.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.167.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.168.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.168.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.168.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.169.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.169.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.169.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.170.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.170.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.170.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.171.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.171.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.171.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.172.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.172.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.172.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.173.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.173.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.173.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.174.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.174.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.174.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.175.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.175.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.175.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.176.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.176.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.176.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.177.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.177.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.177.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.178.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.178.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.178.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.179.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.179.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.179.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.180.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.180.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.180.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.181.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.181.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.181.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.182.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.182.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.182.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.183.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.183.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.183.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.184.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.184.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.184.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.185.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.185.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.185.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.186.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.186.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.186.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.187.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.187.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.187.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.188.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.188.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.188.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.189.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.189.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.189.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.190.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.190.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.190.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.191.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.191.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.191.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.192.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.192.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.192.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.193.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.193.up_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.193.down_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.194.gate_proj.weight": "model-00008-of-000163.safetensors", + "model.layers.5.mlp.experts.194.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.194.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.195.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.195.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.195.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.196.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.196.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.196.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.197.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.197.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.197.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.198.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.198.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.198.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.199.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.199.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.199.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.200.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.200.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.200.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.201.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.201.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.201.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.202.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.202.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.202.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.203.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.203.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.203.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.204.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.204.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.204.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.205.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.205.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.205.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.206.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.206.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.206.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.207.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.207.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.207.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.208.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.208.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.208.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.209.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.209.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.209.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.210.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.210.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.210.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.211.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.211.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.211.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.212.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.212.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.212.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.213.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.213.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.213.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.214.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.214.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.214.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.215.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.215.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.215.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.216.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.216.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.216.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.217.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.217.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.217.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.218.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.218.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.218.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.219.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.219.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.219.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.220.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.220.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.220.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.221.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.221.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.221.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.222.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.222.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.222.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.223.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.223.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.223.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.224.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.224.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.224.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.225.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.225.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.225.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.226.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.226.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.226.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.227.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.227.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.227.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.228.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.228.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.228.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.229.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.229.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.229.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.230.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.230.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.230.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.231.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.231.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.231.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.232.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.232.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.232.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.233.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.233.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.233.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.234.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.234.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.234.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.235.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.235.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.235.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.236.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.236.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.236.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.237.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.237.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.237.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.238.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.238.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.238.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.239.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.239.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.239.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.240.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.240.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.240.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.241.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.241.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.241.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.242.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.242.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.242.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.243.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.243.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.243.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.244.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.244.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.244.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.245.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.245.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.245.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.246.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.246.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.246.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.247.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.247.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.247.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.248.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.248.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.248.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.249.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.249.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.249.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.250.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.250.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.250.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.251.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.251.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.251.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.252.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.252.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.252.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.253.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.253.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.253.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.254.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.254.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.254.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.255.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.255.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.mlp.experts.255.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.5.input_layernorm.weight": "model-00009-of-000163.safetensors", + "model.layers.5.post_attention_layernorm.weight": "model-00009-of-000163.safetensors", + "model.layers.6.self_attn.q_a_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.self_attn.q_a_layernorm.weight": "model-00009-of-000163.safetensors", + "model.layers.6.self_attn.q_b_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.self_attn.kv_a_proj_with_mqa.weight": "model-00009-of-000163.safetensors", + "model.layers.6.self_attn.kv_a_layernorm.weight": "model-00009-of-000163.safetensors", + "model.layers.6.self_attn.kv_b_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.self_attn.o_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.gate.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.gate.e_score_correction_bias": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.shared_experts.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.shared_experts.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.shared_experts.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.0.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.0.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.0.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.1.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.1.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.1.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.2.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.2.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.2.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.3.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.3.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.3.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.4.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.4.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.4.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.5.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.5.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.5.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.6.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.6.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.6.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.7.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.7.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.7.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.8.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.8.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.8.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.9.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.9.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.9.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.10.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.10.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.10.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.11.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.11.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.11.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.12.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.12.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.12.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.13.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.13.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.13.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.14.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.14.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.14.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.15.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.15.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.15.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.16.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.16.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.16.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.17.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.17.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.17.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.18.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.18.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.18.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.19.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.19.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.19.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.20.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.20.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.20.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.21.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.21.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.21.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.22.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.22.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.22.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.23.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.23.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.23.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.24.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.24.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.24.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.25.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.25.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.25.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.26.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.26.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.26.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.27.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.27.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.27.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.28.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.28.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.28.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.29.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.29.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.29.down_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.30.gate_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.30.up_proj.weight": "model-00009-of-000163.safetensors", + "model.layers.6.mlp.experts.30.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.31.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.31.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.31.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.32.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.32.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.32.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.33.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.33.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.33.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.34.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.34.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.34.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.35.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.35.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.35.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.36.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.36.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.36.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.37.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.37.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.37.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.38.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.38.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.38.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.39.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.39.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.39.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.40.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.40.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.40.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.41.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.41.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.41.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.42.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.42.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.42.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.43.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.43.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.43.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.44.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.44.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.44.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.45.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.45.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.45.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.46.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.46.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.46.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.47.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.47.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.47.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.48.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.48.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.48.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.49.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.49.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.49.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.50.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.50.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.50.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.51.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.51.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.51.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.52.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.52.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.52.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.53.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.53.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.53.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.54.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.54.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.54.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.55.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.55.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.55.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.56.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.56.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.56.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.57.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.57.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.57.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.58.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.58.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.58.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.59.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.59.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.59.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.60.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.60.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.60.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.61.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.61.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.61.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.62.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.62.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.62.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.63.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.63.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.63.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.64.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.64.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.64.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.65.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.65.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.65.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.66.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.66.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.66.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.67.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.67.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.67.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.68.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.68.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.68.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.69.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.69.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.69.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.70.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.70.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.70.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.71.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.71.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.71.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.72.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.72.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.72.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.73.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.73.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.73.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.74.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.74.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.74.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.75.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.75.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.75.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.76.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.76.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.76.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.77.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.77.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.77.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.78.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.78.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.78.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.79.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.79.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.79.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.80.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.80.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.80.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.81.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.81.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.81.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.82.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.82.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.82.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.83.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.83.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.83.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.84.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.84.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.84.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.85.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.85.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.85.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.86.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.86.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.86.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.87.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.87.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.87.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.88.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.88.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.88.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.89.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.89.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.89.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.90.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.90.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.90.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.91.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.91.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.91.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.92.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.92.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.92.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.93.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.93.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.93.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.94.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.94.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.94.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.95.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.95.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.95.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.96.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.96.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.96.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.97.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.97.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.97.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.98.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.98.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.98.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.99.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.99.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.99.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.100.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.100.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.100.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.101.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.101.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.101.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.102.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.102.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.102.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.103.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.103.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.103.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.104.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.104.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.104.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.105.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.105.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.105.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.106.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.106.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.106.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.107.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.107.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.107.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.108.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.108.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.108.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.109.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.109.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.109.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.110.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.110.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.110.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.111.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.111.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.111.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.112.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.112.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.112.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.113.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.113.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.113.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.114.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.114.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.114.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.115.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.115.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.115.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.116.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.116.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.116.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.117.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.117.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.117.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.118.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.118.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.118.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.119.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.119.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.119.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.120.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.120.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.120.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.121.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.121.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.121.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.122.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.122.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.122.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.123.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.123.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.123.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.124.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.124.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.124.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.125.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.125.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.125.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.126.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.126.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.126.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.127.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.127.up_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.127.down_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.128.gate_proj.weight": "model-00010-of-000163.safetensors", + "model.layers.6.mlp.experts.128.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.128.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.129.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.129.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.129.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.130.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.130.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.130.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.131.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.131.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.131.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.132.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.132.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.132.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.133.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.133.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.133.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.134.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.134.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.134.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.135.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.135.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.135.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.136.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.136.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.136.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.137.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.137.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.137.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.138.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.138.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.138.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.139.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.139.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.139.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.140.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.140.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.140.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.141.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.141.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.141.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.142.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.142.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.142.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.143.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.143.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.143.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.144.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.144.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.144.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.145.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.145.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.145.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.146.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.146.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.146.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.147.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.147.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.147.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.148.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.148.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.148.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.149.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.149.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.149.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.150.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.150.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.150.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.151.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.151.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.151.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.152.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.152.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.152.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.153.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.153.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.153.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.154.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.154.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.154.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.155.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.155.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.155.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.156.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.156.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.156.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.157.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.157.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.157.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.158.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.158.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.158.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.159.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.159.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.159.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.160.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.160.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.160.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.161.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.161.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.161.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.162.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.162.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.162.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.163.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.163.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.163.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.164.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.164.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.164.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.165.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.165.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.165.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.166.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.166.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.166.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.167.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.167.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.167.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.168.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.168.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.168.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.169.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.169.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.169.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.170.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.170.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.170.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.171.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.171.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.171.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.172.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.172.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.172.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.173.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.173.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.173.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.174.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.174.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.174.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.175.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.175.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.175.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.176.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.176.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.176.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.177.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.177.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.177.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.178.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.178.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.178.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.179.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.179.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.179.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.180.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.180.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.180.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.181.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.181.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.181.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.182.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.182.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.182.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.183.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.183.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.183.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.184.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.184.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.184.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.185.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.185.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.185.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.186.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.186.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.186.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.187.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.187.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.187.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.188.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.188.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.188.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.189.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.189.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.189.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.190.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.190.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.190.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.191.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.191.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.191.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.192.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.192.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.192.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.193.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.193.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.193.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.194.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.194.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.194.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.195.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.195.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.195.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.196.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.196.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.196.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.197.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.197.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.197.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.198.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.198.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.198.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.199.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.199.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.199.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.200.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.200.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.200.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.201.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.201.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.201.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.202.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.202.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.202.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.203.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.203.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.203.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.204.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.204.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.204.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.205.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.205.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.205.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.206.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.206.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.206.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.207.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.207.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.207.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.208.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.208.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.208.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.209.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.209.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.209.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.210.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.210.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.210.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.211.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.211.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.211.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.212.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.212.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.212.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.213.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.213.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.213.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.214.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.214.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.214.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.215.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.215.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.215.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.216.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.216.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.216.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.217.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.217.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.217.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.218.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.218.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.218.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.219.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.219.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.219.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.220.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.220.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.220.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.221.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.221.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.221.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.222.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.222.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.222.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.223.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.223.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.223.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.224.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.224.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.224.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.225.gate_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.225.up_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.225.down_proj.weight": "model-00011-of-000163.safetensors", + "model.layers.6.mlp.experts.226.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.226.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.226.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.227.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.227.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.227.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.228.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.228.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.228.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.229.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.229.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.229.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.230.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.230.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.230.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.231.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.231.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.231.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.232.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.232.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.232.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.233.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.233.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.233.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.234.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.234.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.234.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.235.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.235.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.235.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.236.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.236.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.236.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.237.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.237.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.237.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.238.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.238.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.238.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.239.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.239.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.239.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.240.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.240.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.240.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.241.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.241.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.241.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.242.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.242.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.242.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.243.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.243.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.243.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.244.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.244.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.244.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.245.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.245.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.245.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.246.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.246.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.246.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.247.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.247.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.247.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.248.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.248.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.248.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.249.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.249.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.249.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.250.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.250.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.250.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.251.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.251.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.251.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.252.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.252.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.252.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.253.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.253.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.253.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.254.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.254.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.254.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.255.gate_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.255.up_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.mlp.experts.255.down_proj.weight": "model-00012-of-000163.safetensors", + "model.layers.6.input_layernorm.weight": "model-00012-of-000163.safetensors", + "model.layers.6.post_attention_layernorm.weight": "model-00012-of-000163.safetensors", + "model.layers.7.self_attn.q_a_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.self_attn.q_a_layernorm.weight": "model-00013-of-000163.safetensors", + "model.layers.7.self_attn.q_b_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.self_attn.kv_a_proj_with_mqa.weight": "model-00013-of-000163.safetensors", + "model.layers.7.self_attn.kv_a_layernorm.weight": "model-00013-of-000163.safetensors", + "model.layers.7.self_attn.kv_b_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.self_attn.o_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.gate.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.gate.e_score_correction_bias": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.shared_experts.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.shared_experts.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.shared_experts.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.0.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.0.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.0.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.1.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.1.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.1.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.2.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.2.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.2.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.3.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.3.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.3.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.4.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.4.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.4.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.5.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.5.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.5.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.6.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.6.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.6.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.7.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.7.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.7.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.8.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.8.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.8.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.9.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.9.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.9.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.10.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.10.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.10.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.11.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.11.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.11.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.12.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.12.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.12.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.13.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.13.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.13.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.14.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.14.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.14.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.15.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.15.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.15.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.16.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.16.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.16.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.17.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.17.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.17.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.18.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.18.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.18.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.19.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.19.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.19.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.20.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.20.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.20.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.21.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.21.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.21.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.22.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.22.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.22.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.23.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.23.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.23.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.24.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.24.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.24.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.25.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.25.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.25.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.26.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.26.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.26.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.27.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.27.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.27.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.28.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.28.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.28.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.29.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.29.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.29.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.30.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.30.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.30.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.31.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.31.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.31.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.32.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.32.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.32.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.33.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.33.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.33.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.34.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.34.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.34.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.35.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.35.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.35.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.36.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.36.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.36.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.37.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.37.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.37.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.38.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.38.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.38.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.39.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.39.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.39.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.40.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.40.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.40.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.41.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.41.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.41.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.42.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.42.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.42.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.43.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.43.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.43.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.44.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.44.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.44.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.45.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.45.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.45.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.46.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.46.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.46.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.47.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.47.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.47.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.48.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.48.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.48.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.49.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.49.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.49.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.50.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.50.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.50.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.51.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.51.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.51.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.52.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.52.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.52.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.53.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.53.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.53.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.54.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.54.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.54.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.55.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.55.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.55.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.56.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.56.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.56.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.57.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.57.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.57.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.58.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.58.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.58.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.59.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.59.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.59.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.60.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.60.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.60.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.61.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.61.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.61.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.62.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.62.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.62.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.63.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.63.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.63.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.64.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.64.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.64.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.65.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.65.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.65.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.66.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.66.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.66.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.67.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.67.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.67.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.68.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.68.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.68.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.69.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.69.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.69.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.70.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.70.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.70.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.71.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.71.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.71.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.72.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.72.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.72.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.73.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.73.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.73.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.74.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.74.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.74.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.75.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.75.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.75.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.76.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.76.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.76.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.77.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.77.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.77.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.78.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.78.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.78.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.79.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.79.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.79.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.80.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.80.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.80.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.81.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.81.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.81.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.82.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.82.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.82.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.83.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.83.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.83.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.84.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.84.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.84.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.85.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.85.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.85.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.86.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.86.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.86.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.87.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.87.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.87.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.88.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.88.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.88.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.89.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.89.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.89.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.90.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.90.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.90.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.91.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.91.up_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.91.down_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.92.gate_proj.weight": "model-00013-of-000163.safetensors", + "model.layers.7.mlp.experts.92.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.92.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.93.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.93.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.93.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.94.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.94.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.94.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.95.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.95.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.95.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.96.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.96.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.96.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.97.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.97.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.97.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.98.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.98.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.98.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.99.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.99.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.99.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.100.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.100.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.100.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.101.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.101.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.101.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.102.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.102.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.102.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.103.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.103.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.103.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.104.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.104.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.104.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.105.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.105.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.105.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.106.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.106.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.106.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.107.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.107.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.107.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.108.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.108.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.108.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.109.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.109.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.109.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.110.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.110.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.110.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.111.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.111.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.111.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.112.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.112.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.112.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.113.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.113.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.113.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.114.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.114.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.114.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.115.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.115.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.115.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.116.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.116.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.116.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.117.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.117.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.117.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.118.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.118.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.118.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.119.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.119.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.119.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.120.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.120.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.120.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.121.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.121.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.121.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.122.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.122.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.122.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.123.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.123.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.123.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.124.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.124.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.124.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.125.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.125.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.125.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.126.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.126.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.126.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.127.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.127.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.127.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.128.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.128.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.128.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.129.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.129.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.129.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.130.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.130.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.130.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.131.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.131.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.131.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.132.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.132.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.132.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.133.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.133.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.133.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.134.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.134.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.134.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.135.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.135.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.135.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.136.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.136.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.136.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.137.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.137.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.137.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.138.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.138.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.138.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.139.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.139.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.139.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.140.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.140.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.140.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.141.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.141.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.141.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.142.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.142.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.142.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.143.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.143.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.143.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.144.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.144.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.144.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.145.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.145.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.145.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.146.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.146.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.146.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.147.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.147.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.147.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.148.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.148.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.148.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.149.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.149.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.149.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.150.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.150.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.150.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.151.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.151.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.151.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.152.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.152.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.152.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.153.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.153.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.153.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.154.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.154.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.154.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.155.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.155.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.155.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.156.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.156.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.156.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.157.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.157.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.157.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.158.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.158.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.158.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.159.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.159.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.159.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.160.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.160.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.160.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.161.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.161.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.161.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.162.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.162.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.162.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.163.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.163.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.163.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.164.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.164.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.164.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.165.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.165.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.165.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.166.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.166.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.166.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.167.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.167.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.167.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.168.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.168.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.168.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.169.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.169.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.169.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.170.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.170.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.170.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.171.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.171.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.171.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.172.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.172.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.172.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.173.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.173.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.173.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.174.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.174.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.174.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.175.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.175.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.175.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.176.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.176.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.176.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.177.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.177.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.177.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.178.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.178.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.178.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.179.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.179.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.179.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.180.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.180.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.180.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.181.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.181.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.181.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.182.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.182.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.182.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.183.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.183.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.183.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.184.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.184.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.184.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.185.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.185.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.185.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.186.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.186.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.186.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.187.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.187.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.187.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.188.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.188.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.188.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.189.gate_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.189.up_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.189.down_proj.weight": "model-00014-of-000163.safetensors", + "model.layers.7.mlp.experts.190.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.190.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.190.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.191.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.191.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.191.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.192.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.192.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.192.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.193.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.193.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.193.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.194.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.194.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.194.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.195.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.195.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.195.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.196.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.196.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.196.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.197.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.197.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.197.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.198.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.198.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.198.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.199.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.199.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.199.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.200.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.200.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.200.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.201.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.201.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.201.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.202.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.202.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.202.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.203.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.203.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.203.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.204.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.204.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.204.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.205.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.205.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.205.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.206.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.206.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.206.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.207.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.207.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.207.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.208.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.208.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.208.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.209.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.209.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.209.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.210.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.210.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.210.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.211.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.211.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.211.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.212.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.212.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.212.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.213.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.213.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.213.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.214.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.214.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.214.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.215.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.215.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.215.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.216.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.216.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.216.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.217.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.217.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.217.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.218.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.218.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.218.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.219.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.219.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.219.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.220.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.220.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.220.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.221.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.221.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.221.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.222.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.222.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.222.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.223.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.223.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.223.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.224.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.224.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.224.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.225.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.225.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.225.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.226.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.226.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.226.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.227.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.227.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.227.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.228.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.228.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.228.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.229.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.229.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.229.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.230.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.230.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.230.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.231.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.231.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.231.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.232.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.232.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.232.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.233.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.233.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.233.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.234.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.234.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.234.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.235.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.235.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.235.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.236.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.236.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.236.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.237.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.237.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.237.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.238.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.238.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.238.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.239.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.239.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.239.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.240.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.240.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.240.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.241.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.241.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.241.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.242.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.242.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.242.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.243.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.243.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.243.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.244.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.244.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.244.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.245.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.245.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.245.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.246.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.246.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.246.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.247.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.247.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.247.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.248.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.248.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.248.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.249.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.249.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.249.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.250.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.250.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.250.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.251.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.251.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.251.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.252.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.252.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.252.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.253.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.253.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.253.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.254.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.254.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.254.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.255.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.255.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.mlp.experts.255.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.7.input_layernorm.weight": "model-00015-of-000163.safetensors", + "model.layers.7.post_attention_layernorm.weight": "model-00015-of-000163.safetensors", + "model.layers.8.self_attn.q_a_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.self_attn.q_a_layernorm.weight": "model-00015-of-000163.safetensors", + "model.layers.8.self_attn.q_b_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.self_attn.kv_a_proj_with_mqa.weight": "model-00015-of-000163.safetensors", + "model.layers.8.self_attn.kv_a_layernorm.weight": "model-00015-of-000163.safetensors", + "model.layers.8.self_attn.kv_b_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.self_attn.o_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.gate.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.gate.e_score_correction_bias": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.shared_experts.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.shared_experts.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.shared_experts.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.0.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.0.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.0.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.1.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.1.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.1.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.2.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.2.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.2.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.3.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.3.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.3.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.4.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.4.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.4.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.5.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.5.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.5.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.6.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.6.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.6.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.7.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.7.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.7.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.8.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.8.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.8.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.9.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.9.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.9.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.10.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.10.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.10.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.11.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.11.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.11.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.12.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.12.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.12.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.13.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.13.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.13.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.14.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.14.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.14.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.15.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.15.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.15.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.16.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.16.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.16.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.17.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.17.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.17.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.18.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.18.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.18.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.19.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.19.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.19.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.20.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.20.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.20.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.21.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.21.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.21.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.22.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.22.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.22.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.23.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.23.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.23.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.24.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.24.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.24.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.25.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.25.up_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.25.down_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.26.gate_proj.weight": "model-00015-of-000163.safetensors", + "model.layers.8.mlp.experts.26.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.26.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.27.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.27.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.27.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.28.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.28.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.28.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.29.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.29.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.29.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.30.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.30.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.30.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.31.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.31.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.31.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.32.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.32.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.32.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.33.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.33.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.33.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.34.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.34.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.34.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.35.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.35.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.35.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.36.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.36.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.36.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.37.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.37.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.37.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.38.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.38.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.38.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.39.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.39.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.39.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.40.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.40.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.40.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.41.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.41.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.41.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.42.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.42.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.42.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.43.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.43.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.43.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.44.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.44.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.44.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.45.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.45.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.45.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.46.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.46.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.46.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.47.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.47.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.47.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.48.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.48.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.48.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.49.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.49.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.49.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.50.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.50.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.50.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.51.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.51.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.51.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.52.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.52.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.52.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.53.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.53.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.53.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.54.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.54.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.54.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.55.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.55.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.55.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.56.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.56.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.56.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.57.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.57.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.57.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.58.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.58.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.58.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.59.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.59.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.59.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.60.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.60.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.60.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.61.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.61.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.61.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.62.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.62.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.62.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.63.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.63.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.63.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.64.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.64.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.64.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.65.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.65.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.65.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.66.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.66.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.66.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.67.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.67.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.67.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.68.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.68.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.68.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.69.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.69.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.69.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.70.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.70.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.70.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.71.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.71.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.71.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.72.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.72.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.72.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.73.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.73.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.73.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.74.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.74.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.74.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.75.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.75.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.75.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.76.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.76.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.76.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.77.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.77.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.77.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.78.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.78.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.78.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.79.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.79.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.79.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.80.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.80.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.80.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.81.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.81.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.81.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.82.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.82.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.82.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.83.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.83.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.83.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.84.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.84.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.84.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.85.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.85.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.85.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.86.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.86.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.86.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.87.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.87.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.87.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.88.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.88.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.88.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.89.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.89.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.89.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.90.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.90.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.90.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.91.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.91.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.91.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.92.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.92.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.92.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.93.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.93.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.93.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.94.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.94.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.94.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.95.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.95.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.95.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.96.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.96.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.96.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.97.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.97.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.97.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.98.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.98.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.98.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.99.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.99.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.99.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.100.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.100.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.100.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.101.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.101.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.101.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.102.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.102.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.102.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.103.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.103.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.103.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.104.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.104.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.104.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.105.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.105.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.105.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.106.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.106.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.106.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.107.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.107.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.107.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.108.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.108.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.108.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.109.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.109.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.109.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.110.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.110.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.110.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.111.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.111.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.111.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.112.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.112.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.112.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.113.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.113.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.113.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.114.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.114.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.114.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.115.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.115.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.115.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.116.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.116.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.116.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.117.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.117.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.117.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.118.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.118.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.118.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.119.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.119.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.119.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.120.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.120.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.120.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.121.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.121.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.121.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.122.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.122.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.122.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.123.gate_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.123.up_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.123.down_proj.weight": "model-00016-of-000163.safetensors", + "model.layers.8.mlp.experts.124.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.124.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.124.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.125.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.125.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.125.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.126.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.126.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.126.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.127.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.127.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.127.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.128.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.128.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.128.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.129.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.129.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.129.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.130.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.130.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.130.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.131.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.131.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.131.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.132.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.132.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.132.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.133.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.133.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.133.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.134.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.134.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.134.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.135.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.135.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.135.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.136.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.136.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.136.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.137.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.137.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.137.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.138.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.138.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.138.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.139.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.139.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.139.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.140.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.140.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.140.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.141.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.141.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.141.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.142.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.142.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.142.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.143.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.143.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.143.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.144.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.144.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.144.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.145.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.145.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.145.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.146.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.146.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.146.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.147.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.147.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.147.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.148.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.148.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.148.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.149.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.149.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.149.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.150.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.150.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.150.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.151.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.151.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.151.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.152.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.152.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.152.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.153.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.153.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.153.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.154.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.154.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.154.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.155.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.155.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.155.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.156.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.156.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.156.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.157.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.157.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.157.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.158.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.158.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.158.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.159.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.159.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.159.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.160.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.160.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.160.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.161.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.161.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.161.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.162.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.162.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.162.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.163.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.163.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.163.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.164.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.164.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.164.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.165.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.165.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.165.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.166.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.166.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.166.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.167.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.167.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.167.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.168.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.168.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.168.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.169.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.169.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.169.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.170.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.170.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.170.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.171.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.171.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.171.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.172.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.172.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.172.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.173.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.173.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.173.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.174.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.174.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.174.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.175.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.175.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.175.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.176.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.176.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.176.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.177.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.177.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.177.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.178.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.178.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.178.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.179.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.179.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.179.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.180.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.180.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.180.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.181.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.181.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.181.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.182.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.182.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.182.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.183.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.183.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.183.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.184.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.184.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.184.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.185.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.185.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.185.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.186.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.186.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.186.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.187.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.187.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.187.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.188.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.188.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.188.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.189.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.189.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.189.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.190.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.190.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.190.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.191.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.191.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.191.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.192.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.192.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.192.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.193.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.193.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.193.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.194.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.194.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.194.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.195.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.195.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.195.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.196.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.196.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.196.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.197.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.197.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.197.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.198.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.198.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.198.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.199.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.199.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.199.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.200.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.200.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.200.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.201.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.201.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.201.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.202.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.202.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.202.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.203.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.203.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.203.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.204.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.204.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.204.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.205.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.205.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.205.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.206.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.206.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.206.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.207.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.207.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.207.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.208.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.208.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.208.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.209.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.209.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.209.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.210.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.210.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.210.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.211.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.211.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.211.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.212.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.212.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.212.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.213.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.213.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.213.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.214.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.214.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.214.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.215.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.215.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.215.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.216.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.216.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.216.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.217.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.217.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.217.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.218.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.218.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.218.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.219.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.219.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.219.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.220.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.220.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.220.down_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.221.gate_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.221.up_proj.weight": "model-00017-of-000163.safetensors", + "model.layers.8.mlp.experts.221.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.222.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.222.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.222.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.223.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.223.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.223.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.224.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.224.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.224.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.225.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.225.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.225.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.226.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.226.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.226.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.227.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.227.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.227.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.228.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.228.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.228.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.229.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.229.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.229.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.230.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.230.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.230.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.231.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.231.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.231.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.232.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.232.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.232.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.233.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.233.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.233.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.234.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.234.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.234.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.235.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.235.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.235.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.236.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.236.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.236.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.237.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.237.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.237.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.238.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.238.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.238.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.239.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.239.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.239.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.240.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.240.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.240.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.241.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.241.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.241.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.242.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.242.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.242.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.243.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.243.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.243.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.244.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.244.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.244.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.245.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.245.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.245.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.246.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.246.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.246.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.247.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.247.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.247.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.248.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.248.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.248.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.249.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.249.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.249.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.250.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.250.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.250.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.251.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.251.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.251.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.252.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.252.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.252.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.253.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.253.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.253.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.254.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.254.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.254.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.255.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.255.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.mlp.experts.255.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.8.input_layernorm.weight": "model-00018-of-000163.safetensors", + "model.layers.8.post_attention_layernorm.weight": "model-00018-of-000163.safetensors", + "model.layers.9.self_attn.q_a_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.self_attn.q_a_layernorm.weight": "model-00018-of-000163.safetensors", + "model.layers.9.self_attn.q_b_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.self_attn.kv_a_proj_with_mqa.weight": "model-00018-of-000163.safetensors", + "model.layers.9.self_attn.kv_a_layernorm.weight": "model-00018-of-000163.safetensors", + "model.layers.9.self_attn.kv_b_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.self_attn.o_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.gate.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.gate.e_score_correction_bias": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.shared_experts.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.shared_experts.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.shared_experts.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.0.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.0.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.0.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.1.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.1.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.1.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.2.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.2.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.2.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.3.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.3.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.3.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.4.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.4.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.4.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.5.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.5.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.5.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.6.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.6.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.6.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.7.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.7.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.7.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.8.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.8.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.8.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.9.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.9.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.9.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.10.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.10.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.10.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.11.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.11.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.11.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.12.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.12.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.12.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.13.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.13.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.13.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.14.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.14.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.14.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.15.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.15.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.15.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.16.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.16.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.16.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.17.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.17.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.17.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.18.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.18.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.18.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.19.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.19.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.19.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.20.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.20.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.20.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.21.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.21.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.21.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.22.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.22.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.22.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.23.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.23.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.23.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.24.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.24.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.24.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.25.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.25.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.25.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.26.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.26.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.26.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.27.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.27.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.27.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.28.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.28.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.28.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.29.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.29.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.29.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.30.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.30.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.30.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.31.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.31.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.31.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.32.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.32.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.32.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.33.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.33.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.33.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.34.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.34.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.34.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.35.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.35.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.35.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.36.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.36.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.36.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.37.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.37.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.37.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.38.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.38.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.38.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.39.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.39.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.39.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.40.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.40.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.40.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.41.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.41.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.41.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.42.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.42.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.42.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.43.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.43.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.43.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.44.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.44.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.44.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.45.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.45.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.45.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.46.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.46.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.46.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.47.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.47.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.47.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.48.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.48.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.48.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.49.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.49.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.49.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.50.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.50.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.50.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.51.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.51.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.51.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.52.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.52.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.52.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.53.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.53.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.53.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.54.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.54.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.54.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.55.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.55.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.55.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.56.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.56.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.56.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.57.gate_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.57.up_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.57.down_proj.weight": "model-00018-of-000163.safetensors", + "model.layers.9.mlp.experts.58.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.58.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.58.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.59.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.59.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.59.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.60.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.60.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.60.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.61.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.61.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.61.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.62.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.62.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.62.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.63.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.63.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.63.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.64.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.64.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.64.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.65.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.65.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.65.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.66.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.66.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.66.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.67.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.67.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.67.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.68.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.68.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.68.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.69.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.69.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.69.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.70.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.70.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.70.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.71.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.71.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.71.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.72.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.72.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.72.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.73.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.73.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.73.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.74.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.74.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.74.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.75.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.75.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.75.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.76.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.76.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.76.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.77.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.77.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.77.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.78.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.78.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.78.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.79.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.79.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.79.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.80.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.80.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.80.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.81.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.81.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.81.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.82.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.82.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.82.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.83.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.83.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.83.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.84.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.84.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.84.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.85.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.85.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.85.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.86.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.86.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.86.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.87.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.87.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.87.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.88.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.88.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.88.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.89.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.89.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.89.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.90.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.90.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.90.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.91.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.91.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.91.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.92.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.92.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.92.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.93.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.93.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.93.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.94.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.94.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.94.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.95.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.95.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.95.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.96.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.96.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.96.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.97.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.97.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.97.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.98.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.98.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.98.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.99.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.99.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.99.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.100.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.100.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.100.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.101.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.101.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.101.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.102.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.102.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.102.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.103.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.103.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.103.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.104.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.104.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.104.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.105.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.105.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.105.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.106.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.106.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.106.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.107.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.107.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.107.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.108.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.108.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.108.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.109.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.109.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.109.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.110.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.110.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.110.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.111.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.111.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.111.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.112.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.112.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.112.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.113.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.113.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.113.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.114.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.114.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.114.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.115.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.115.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.115.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.116.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.116.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.116.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.117.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.117.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.117.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.118.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.118.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.118.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.119.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.119.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.119.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.120.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.120.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.120.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.121.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.121.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.121.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.122.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.122.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.122.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.123.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.123.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.123.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.124.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.124.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.124.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.125.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.125.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.125.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.126.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.126.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.126.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.127.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.127.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.127.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.128.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.128.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.128.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.129.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.129.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.129.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.130.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.130.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.130.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.131.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.131.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.131.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.132.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.132.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.132.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.133.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.133.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.133.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.134.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.134.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.134.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.135.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.135.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.135.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.136.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.136.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.136.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.137.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.137.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.137.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.138.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.138.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.138.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.139.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.139.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.139.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.140.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.140.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.140.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.141.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.141.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.141.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.142.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.142.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.142.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.143.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.143.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.143.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.144.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.144.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.144.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.145.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.145.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.145.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.146.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.146.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.146.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.147.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.147.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.147.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.148.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.148.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.148.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.149.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.149.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.149.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.150.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.150.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.150.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.151.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.151.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.151.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.152.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.152.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.152.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.153.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.153.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.153.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.154.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.154.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.154.down_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.155.gate_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.155.up_proj.weight": "model-00019-of-000163.safetensors", + "model.layers.9.mlp.experts.155.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.156.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.156.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.156.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.157.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.157.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.157.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.158.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.158.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.158.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.159.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.159.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.159.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.160.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.160.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.160.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.161.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.161.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.161.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.162.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.162.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.162.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.163.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.163.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.163.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.164.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.164.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.164.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.165.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.165.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.165.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.166.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.166.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.166.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.167.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.167.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.167.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.168.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.168.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.168.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.169.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.169.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.169.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.170.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.170.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.170.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.171.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.171.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.171.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.172.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.172.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.172.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.173.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.173.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.173.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.174.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.174.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.174.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.175.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.175.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.175.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.176.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.176.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.176.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.177.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.177.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.177.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.178.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.178.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.178.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.179.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.179.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.179.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.180.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.180.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.180.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.181.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.181.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.181.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.182.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.182.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.182.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.183.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.183.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.183.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.184.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.184.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.184.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.185.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.185.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.185.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.186.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.186.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.186.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.187.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.187.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.187.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.188.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.188.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.188.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.189.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.189.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.189.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.190.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.190.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.190.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.191.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.191.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.191.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.192.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.192.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.192.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.193.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.193.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.193.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.194.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.194.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.194.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.195.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.195.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.195.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.196.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.196.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.196.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.197.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.197.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.197.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.198.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.198.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.198.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.199.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.199.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.199.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.200.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.200.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.200.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.201.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.201.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.201.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.202.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.202.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.202.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.203.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.203.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.203.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.204.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.204.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.204.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.205.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.205.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.205.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.206.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.206.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.206.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.207.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.207.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.207.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.208.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.208.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.208.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.209.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.209.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.209.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.210.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.210.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.210.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.211.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.211.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.211.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.212.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.212.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.212.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.213.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.213.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.213.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.214.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.214.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.214.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.215.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.215.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.215.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.216.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.216.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.216.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.217.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.217.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.217.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.218.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.218.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.218.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.219.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.219.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.219.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.220.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.220.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.220.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.221.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.221.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.221.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.222.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.222.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.222.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.223.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.223.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.223.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.224.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.224.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.224.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.225.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.225.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.225.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.226.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.226.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.226.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.227.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.227.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.227.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.228.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.228.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.228.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.229.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.229.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.229.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.230.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.230.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.230.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.231.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.231.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.231.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.232.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.232.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.232.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.233.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.233.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.233.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.234.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.234.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.234.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.235.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.235.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.235.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.236.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.236.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.236.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.237.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.237.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.237.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.238.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.238.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.238.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.239.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.239.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.239.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.240.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.240.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.240.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.241.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.241.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.241.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.242.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.242.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.242.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.243.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.243.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.243.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.244.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.244.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.244.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.245.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.245.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.245.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.246.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.246.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.246.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.247.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.247.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.247.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.248.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.248.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.248.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.249.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.249.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.249.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.250.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.250.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.250.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.251.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.251.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.251.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.252.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.252.up_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.252.down_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.253.gate_proj.weight": "model-00020-of-000163.safetensors", + "model.layers.9.mlp.experts.253.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.9.mlp.experts.253.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.9.mlp.experts.254.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.9.mlp.experts.254.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.9.mlp.experts.254.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.9.mlp.experts.255.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.9.mlp.experts.255.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.9.mlp.experts.255.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.9.input_layernorm.weight": "model-00021-of-000163.safetensors", + "model.layers.9.post_attention_layernorm.weight": "model-00021-of-000163.safetensors", + "model.layers.10.self_attn.q_a_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.self_attn.q_a_layernorm.weight": "model-00021-of-000163.safetensors", + "model.layers.10.self_attn.q_b_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.self_attn.kv_a_proj_with_mqa.weight": "model-00021-of-000163.safetensors", + "model.layers.10.self_attn.kv_a_layernorm.weight": "model-00021-of-000163.safetensors", + "model.layers.10.self_attn.kv_b_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.self_attn.o_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.gate.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.gate.e_score_correction_bias": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.shared_experts.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.shared_experts.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.shared_experts.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.0.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.0.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.0.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.1.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.1.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.1.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.2.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.2.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.2.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.3.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.3.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.3.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.4.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.4.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.4.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.5.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.5.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.5.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.6.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.6.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.6.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.7.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.7.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.7.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.8.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.8.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.8.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.9.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.9.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.9.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.10.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.10.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.10.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.11.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.11.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.11.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.12.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.12.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.12.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.13.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.13.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.13.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.14.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.14.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.14.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.15.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.15.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.15.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.16.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.16.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.16.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.17.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.17.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.17.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.18.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.18.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.18.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.19.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.19.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.19.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.20.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.20.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.20.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.21.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.21.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.21.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.22.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.22.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.22.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.23.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.23.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.23.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.24.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.24.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.24.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.25.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.25.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.25.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.26.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.26.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.26.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.27.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.27.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.27.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.28.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.28.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.28.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.29.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.29.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.29.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.30.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.30.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.30.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.31.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.31.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.31.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.32.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.32.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.32.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.33.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.33.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.33.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.34.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.34.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.34.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.35.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.35.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.35.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.36.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.36.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.36.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.37.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.37.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.37.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.38.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.38.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.38.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.39.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.39.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.39.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.40.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.40.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.40.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.41.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.41.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.41.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.42.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.42.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.42.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.43.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.43.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.43.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.44.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.44.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.44.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.45.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.45.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.45.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.46.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.46.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.46.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.47.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.47.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.47.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.48.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.48.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.48.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.49.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.49.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.49.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.50.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.50.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.50.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.51.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.51.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.51.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.52.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.52.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.52.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.53.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.53.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.53.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.54.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.54.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.54.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.55.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.55.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.55.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.56.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.56.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.56.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.57.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.57.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.57.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.58.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.58.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.58.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.59.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.59.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.59.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.60.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.60.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.60.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.61.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.61.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.61.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.62.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.62.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.62.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.63.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.63.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.63.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.64.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.64.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.64.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.65.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.65.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.65.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.66.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.66.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.66.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.67.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.67.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.67.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.68.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.68.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.68.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.69.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.69.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.69.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.70.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.70.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.70.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.71.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.71.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.71.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.72.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.72.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.72.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.73.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.73.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.73.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.74.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.74.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.74.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.75.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.75.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.75.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.76.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.76.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.76.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.77.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.77.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.77.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.78.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.78.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.78.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.79.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.79.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.79.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.80.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.80.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.80.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.81.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.81.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.81.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.82.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.82.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.82.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.83.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.83.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.83.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.84.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.84.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.84.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.85.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.85.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.85.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.86.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.86.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.86.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.87.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.87.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.87.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.88.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.88.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.88.down_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.89.gate_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.89.up_proj.weight": "model-00021-of-000163.safetensors", + "model.layers.10.mlp.experts.89.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.90.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.90.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.90.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.91.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.91.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.91.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.92.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.92.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.92.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.93.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.93.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.93.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.94.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.94.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.94.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.95.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.95.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.95.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.96.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.96.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.96.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.97.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.97.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.97.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.98.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.98.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.98.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.99.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.99.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.99.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.100.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.100.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.100.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.101.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.101.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.101.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.102.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.102.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.102.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.103.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.103.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.103.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.104.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.104.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.104.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.105.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.105.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.105.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.106.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.106.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.106.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.107.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.107.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.107.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.108.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.108.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.108.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.109.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.109.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.109.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.110.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.110.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.110.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.111.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.111.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.111.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.112.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.112.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.112.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.113.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.113.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.113.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.114.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.114.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.114.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.115.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.115.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.115.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.116.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.116.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.116.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.117.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.117.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.117.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.118.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.118.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.118.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.119.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.119.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.119.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.120.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.120.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.120.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.121.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.121.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.121.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.122.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.122.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.122.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.123.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.123.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.123.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.124.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.124.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.124.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.125.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.125.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.125.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.126.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.126.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.126.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.127.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.127.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.127.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.128.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.128.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.128.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.129.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.129.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.129.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.130.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.130.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.130.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.131.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.131.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.131.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.132.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.132.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.132.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.133.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.133.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.133.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.134.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.134.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.134.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.135.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.135.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.135.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.136.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.136.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.136.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.137.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.137.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.137.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.138.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.138.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.138.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.139.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.139.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.139.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.140.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.140.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.140.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.141.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.141.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.141.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.142.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.142.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.142.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.143.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.143.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.143.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.144.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.144.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.144.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.145.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.145.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.145.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.146.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.146.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.146.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.147.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.147.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.147.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.148.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.148.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.148.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.149.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.149.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.149.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.150.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.150.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.150.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.151.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.151.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.151.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.152.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.152.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.152.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.153.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.153.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.153.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.154.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.154.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.154.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.155.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.155.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.155.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.156.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.156.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.156.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.157.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.157.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.157.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.158.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.158.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.158.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.159.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.159.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.159.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.160.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.160.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.160.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.161.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.161.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.161.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.162.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.162.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.162.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.163.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.163.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.163.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.164.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.164.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.164.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.165.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.165.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.165.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.166.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.166.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.166.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.167.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.167.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.167.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.168.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.168.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.168.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.169.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.169.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.169.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.170.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.170.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.170.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.171.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.171.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.171.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.172.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.172.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.172.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.173.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.173.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.173.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.174.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.174.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.174.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.175.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.175.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.175.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.176.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.176.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.176.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.177.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.177.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.177.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.178.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.178.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.178.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.179.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.179.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.179.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.180.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.180.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.180.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.181.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.181.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.181.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.182.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.182.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.182.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.183.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.183.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.183.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.184.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.184.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.184.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.185.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.185.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.185.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.186.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.186.up_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.186.down_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.187.gate_proj.weight": "model-00022-of-000163.safetensors", + "model.layers.10.mlp.experts.187.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.187.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.188.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.188.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.188.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.189.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.189.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.189.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.190.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.190.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.190.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.191.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.191.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.191.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.192.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.192.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.192.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.193.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.193.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.193.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.194.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.194.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.194.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.195.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.195.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.195.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.196.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.196.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.196.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.197.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.197.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.197.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.198.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.198.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.198.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.199.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.199.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.199.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.200.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.200.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.200.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.201.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.201.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.201.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.202.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.202.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.202.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.203.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.203.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.203.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.204.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.204.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.204.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.205.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.205.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.205.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.206.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.206.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.206.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.207.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.207.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.207.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.208.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.208.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.208.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.209.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.209.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.209.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.210.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.210.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.210.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.211.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.211.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.211.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.212.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.212.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.212.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.213.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.213.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.213.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.214.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.214.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.214.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.215.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.215.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.215.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.216.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.216.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.216.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.217.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.217.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.217.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.218.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.218.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.218.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.219.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.219.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.219.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.220.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.220.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.220.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.221.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.221.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.221.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.222.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.222.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.222.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.223.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.223.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.223.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.224.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.224.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.224.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.225.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.225.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.225.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.226.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.226.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.226.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.227.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.227.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.227.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.228.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.228.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.228.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.229.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.229.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.229.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.230.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.230.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.230.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.231.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.231.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.231.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.232.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.232.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.232.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.233.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.233.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.233.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.234.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.234.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.234.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.235.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.235.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.235.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.236.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.236.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.236.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.237.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.237.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.237.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.238.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.238.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.238.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.239.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.239.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.239.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.240.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.240.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.240.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.241.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.241.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.241.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.242.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.242.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.242.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.243.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.243.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.243.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.244.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.244.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.244.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.245.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.245.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.245.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.246.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.246.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.246.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.247.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.247.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.247.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.248.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.248.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.248.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.249.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.249.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.249.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.250.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.250.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.250.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.251.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.251.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.251.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.252.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.252.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.252.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.253.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.253.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.253.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.254.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.254.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.254.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.255.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.255.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.mlp.experts.255.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.10.input_layernorm.weight": "model-00023-of-000163.safetensors", + "model.layers.10.post_attention_layernorm.weight": "model-00023-of-000163.safetensors", + "model.layers.11.self_attn.q_a_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.self_attn.q_a_layernorm.weight": "model-00023-of-000163.safetensors", + "model.layers.11.self_attn.q_b_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.self_attn.kv_a_proj_with_mqa.weight": "model-00023-of-000163.safetensors", + "model.layers.11.self_attn.kv_a_layernorm.weight": "model-00023-of-000163.safetensors", + "model.layers.11.self_attn.kv_b_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.self_attn.o_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.gate.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.gate.e_score_correction_bias": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.shared_experts.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.shared_experts.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.shared_experts.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.0.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.0.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.0.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.1.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.1.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.1.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.2.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.2.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.2.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.3.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.3.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.3.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.4.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.4.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.4.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.5.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.5.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.5.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.6.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.6.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.6.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.7.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.7.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.7.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.8.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.8.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.8.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.9.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.9.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.9.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.10.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.10.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.10.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.11.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.11.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.11.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.12.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.12.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.12.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.13.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.13.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.13.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.14.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.14.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.14.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.15.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.15.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.15.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.16.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.16.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.16.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.17.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.17.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.17.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.18.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.18.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.18.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.19.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.19.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.19.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.20.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.20.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.20.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.21.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.21.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.21.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.22.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.22.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.22.down_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.23.gate_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.23.up_proj.weight": "model-00023-of-000163.safetensors", + "model.layers.11.mlp.experts.23.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.24.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.24.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.24.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.25.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.25.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.25.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.26.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.26.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.26.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.27.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.27.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.27.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.28.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.28.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.28.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.29.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.29.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.29.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.30.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.30.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.30.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.31.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.31.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.31.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.32.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.32.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.32.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.33.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.33.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.33.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.34.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.34.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.34.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.35.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.35.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.35.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.36.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.36.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.36.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.37.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.37.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.37.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.38.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.38.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.38.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.39.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.39.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.39.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.40.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.40.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.40.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.41.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.41.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.41.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.42.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.42.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.42.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.43.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.43.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.43.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.44.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.44.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.44.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.45.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.45.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.45.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.46.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.46.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.46.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.47.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.47.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.47.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.48.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.48.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.48.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.49.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.49.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.49.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.50.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.50.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.50.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.51.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.51.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.51.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.52.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.52.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.52.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.53.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.53.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.53.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.54.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.54.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.54.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.55.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.55.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.55.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.56.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.56.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.56.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.57.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.57.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.57.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.58.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.58.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.58.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.59.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.59.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.59.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.60.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.60.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.60.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.61.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.61.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.61.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.62.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.62.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.62.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.63.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.63.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.63.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.64.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.64.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.64.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.65.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.65.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.65.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.66.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.66.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.66.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.67.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.67.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.67.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.68.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.68.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.68.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.69.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.69.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.69.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.70.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.70.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.70.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.71.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.71.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.71.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.72.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.72.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.72.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.73.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.73.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.73.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.74.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.74.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.74.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.75.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.75.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.75.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.76.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.76.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.76.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.77.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.77.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.77.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.78.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.78.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.78.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.79.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.79.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.79.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.80.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.80.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.80.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.81.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.81.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.81.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.82.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.82.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.82.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.83.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.83.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.83.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.84.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.84.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.84.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.85.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.85.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.85.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.86.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.86.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.86.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.87.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.87.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.87.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.88.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.88.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.88.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.89.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.89.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.89.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.90.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.90.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.90.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.91.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.91.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.91.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.92.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.92.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.92.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.93.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.93.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.93.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.94.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.94.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.94.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.95.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.95.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.95.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.96.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.96.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.96.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.97.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.97.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.97.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.98.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.98.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.98.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.99.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.99.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.99.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.100.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.100.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.100.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.101.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.101.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.101.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.102.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.102.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.102.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.103.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.103.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.103.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.104.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.104.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.104.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.105.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.105.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.105.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.106.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.106.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.106.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.107.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.107.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.107.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.108.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.108.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.108.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.109.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.109.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.109.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.110.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.110.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.110.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.111.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.111.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.111.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.112.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.112.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.112.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.113.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.113.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.113.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.114.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.114.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.114.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.115.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.115.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.115.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.116.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.116.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.116.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.117.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.117.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.117.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.118.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.118.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.118.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.119.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.119.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.119.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.120.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.120.up_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.120.down_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.121.gate_proj.weight": "model-00024-of-000163.safetensors", + "model.layers.11.mlp.experts.121.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.121.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.122.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.122.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.122.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.123.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.123.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.123.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.124.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.124.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.124.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.125.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.125.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.125.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.126.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.126.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.126.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.127.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.127.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.127.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.128.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.128.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.128.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.129.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.129.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.129.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.130.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.130.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.130.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.131.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.131.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.131.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.132.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.132.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.132.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.133.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.133.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.133.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.134.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.134.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.134.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.135.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.135.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.135.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.136.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.136.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.136.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.137.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.137.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.137.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.138.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.138.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.138.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.139.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.139.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.139.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.140.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.140.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.140.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.141.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.141.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.141.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.142.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.142.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.142.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.143.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.143.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.143.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.144.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.144.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.144.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.145.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.145.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.145.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.146.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.146.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.146.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.147.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.147.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.147.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.148.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.148.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.148.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.149.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.149.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.149.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.150.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.150.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.150.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.151.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.151.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.151.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.152.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.152.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.152.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.153.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.153.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.153.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.154.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.154.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.154.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.155.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.155.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.155.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.156.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.156.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.156.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.157.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.157.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.157.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.158.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.158.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.158.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.159.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.159.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.159.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.160.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.160.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.160.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.161.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.161.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.161.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.162.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.162.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.162.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.163.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.163.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.163.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.164.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.164.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.164.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.165.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.165.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.165.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.166.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.166.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.166.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.167.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.167.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.167.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.168.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.168.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.168.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.169.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.169.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.169.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.170.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.170.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.170.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.171.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.171.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.171.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.172.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.172.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.172.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.173.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.173.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.173.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.174.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.174.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.174.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.175.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.175.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.175.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.176.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.176.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.176.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.177.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.177.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.177.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.178.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.178.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.178.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.179.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.179.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.179.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.180.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.180.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.180.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.181.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.181.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.181.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.182.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.182.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.182.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.183.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.183.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.183.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.184.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.184.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.184.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.185.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.185.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.185.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.186.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.186.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.186.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.187.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.187.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.187.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.188.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.188.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.188.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.189.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.189.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.189.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.190.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.190.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.190.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.191.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.191.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.191.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.192.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.192.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.192.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.193.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.193.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.193.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.194.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.194.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.194.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.195.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.195.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.195.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.196.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.196.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.196.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.197.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.197.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.197.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.198.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.198.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.198.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.199.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.199.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.199.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.200.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.200.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.200.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.201.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.201.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.201.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.202.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.202.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.202.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.203.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.203.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.203.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.204.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.204.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.204.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.205.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.205.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.205.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.206.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.206.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.206.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.207.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.207.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.207.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.208.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.208.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.208.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.209.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.209.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.209.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.210.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.210.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.210.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.211.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.211.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.211.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.212.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.212.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.212.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.213.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.213.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.213.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.214.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.214.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.214.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.215.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.215.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.215.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.216.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.216.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.216.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.217.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.217.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.217.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.218.gate_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.218.up_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.218.down_proj.weight": "model-00025-of-000163.safetensors", + "model.layers.11.mlp.experts.219.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.219.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.219.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.220.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.220.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.220.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.221.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.221.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.221.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.222.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.222.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.222.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.223.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.223.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.223.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.224.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.224.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.224.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.225.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.225.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.225.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.226.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.226.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.226.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.227.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.227.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.227.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.228.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.228.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.228.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.229.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.229.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.229.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.230.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.230.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.230.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.231.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.231.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.231.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.232.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.232.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.232.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.233.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.233.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.233.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.234.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.234.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.234.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.235.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.235.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.235.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.236.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.236.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.236.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.237.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.237.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.237.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.238.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.238.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.238.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.239.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.239.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.239.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.240.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.240.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.240.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.241.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.241.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.241.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.242.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.242.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.242.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.243.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.243.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.243.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.244.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.244.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.244.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.245.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.245.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.245.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.246.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.246.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.246.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.247.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.247.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.247.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.248.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.248.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.248.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.249.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.249.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.249.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.250.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.250.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.250.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.251.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.251.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.251.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.252.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.252.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.252.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.253.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.253.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.253.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.254.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.254.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.254.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.255.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.255.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.mlp.experts.255.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.11.input_layernorm.weight": "model-00026-of-000163.safetensors", + "model.layers.11.post_attention_layernorm.weight": "model-00026-of-000163.safetensors", + "model.layers.12.self_attn.q_a_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.self_attn.q_a_layernorm.weight": "model-00026-of-000163.safetensors", + "model.layers.12.self_attn.q_b_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.self_attn.kv_a_proj_with_mqa.weight": "model-00026-of-000163.safetensors", + "model.layers.12.self_attn.kv_a_layernorm.weight": "model-00026-of-000163.safetensors", + "model.layers.12.self_attn.kv_b_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.self_attn.o_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.gate.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.gate.e_score_correction_bias": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.shared_experts.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.shared_experts.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.shared_experts.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.0.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.0.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.0.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.1.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.1.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.1.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.2.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.2.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.2.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.3.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.3.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.3.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.4.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.4.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.4.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.5.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.5.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.5.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.6.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.6.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.6.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.7.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.7.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.7.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.8.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.8.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.8.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.9.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.9.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.9.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.10.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.10.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.10.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.11.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.11.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.11.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.12.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.12.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.12.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.13.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.13.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.13.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.14.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.14.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.14.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.15.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.15.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.15.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.16.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.16.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.16.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.17.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.17.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.17.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.18.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.18.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.18.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.19.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.19.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.19.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.20.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.20.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.20.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.21.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.21.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.21.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.22.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.22.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.22.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.23.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.23.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.23.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.24.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.24.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.24.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.25.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.25.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.25.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.26.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.26.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.26.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.27.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.27.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.27.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.28.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.28.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.28.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.29.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.29.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.29.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.30.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.30.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.30.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.31.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.31.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.31.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.32.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.32.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.32.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.33.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.33.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.33.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.34.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.34.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.34.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.35.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.35.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.35.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.36.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.36.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.36.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.37.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.37.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.37.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.38.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.38.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.38.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.39.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.39.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.39.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.40.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.40.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.40.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.41.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.41.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.41.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.42.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.42.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.42.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.43.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.43.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.43.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.44.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.44.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.44.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.45.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.45.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.45.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.46.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.46.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.46.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.47.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.47.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.47.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.48.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.48.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.48.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.49.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.49.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.49.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.50.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.50.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.50.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.51.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.51.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.51.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.52.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.52.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.52.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.53.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.53.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.53.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.54.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.54.up_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.54.down_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.55.gate_proj.weight": "model-00026-of-000163.safetensors", + "model.layers.12.mlp.experts.55.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.55.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.56.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.56.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.56.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.57.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.57.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.57.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.58.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.58.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.58.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.59.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.59.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.59.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.60.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.60.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.60.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.61.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.61.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.61.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.62.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.62.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.62.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.63.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.63.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.63.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.64.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.64.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.64.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.65.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.65.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.65.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.66.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.66.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.66.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.67.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.67.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.67.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.68.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.68.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.68.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.69.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.69.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.69.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.70.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.70.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.70.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.71.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.71.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.71.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.72.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.72.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.72.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.73.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.73.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.73.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.74.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.74.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.74.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.75.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.75.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.75.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.76.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.76.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.76.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.77.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.77.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.77.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.78.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.78.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.78.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.79.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.79.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.79.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.80.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.80.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.80.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.81.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.81.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.81.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.82.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.82.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.82.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.83.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.83.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.83.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.84.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.84.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.84.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.85.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.85.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.85.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.86.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.86.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.86.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.87.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.87.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.87.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.88.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.88.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.88.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.89.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.89.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.89.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.90.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.90.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.90.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.91.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.91.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.91.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.92.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.92.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.92.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.93.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.93.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.93.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.94.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.94.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.94.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.95.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.95.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.95.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.96.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.96.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.96.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.97.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.97.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.97.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.98.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.98.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.98.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.99.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.99.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.99.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.100.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.100.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.100.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.101.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.101.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.101.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.102.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.102.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.102.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.103.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.103.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.103.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.104.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.104.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.104.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.105.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.105.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.105.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.106.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.106.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.106.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.107.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.107.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.107.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.108.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.108.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.108.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.109.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.109.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.109.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.110.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.110.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.110.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.111.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.111.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.111.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.112.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.112.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.112.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.113.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.113.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.113.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.114.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.114.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.114.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.115.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.115.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.115.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.116.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.116.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.116.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.117.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.117.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.117.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.118.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.118.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.118.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.119.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.119.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.119.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.120.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.120.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.120.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.121.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.121.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.121.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.122.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.122.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.122.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.123.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.123.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.123.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.124.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.124.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.124.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.125.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.125.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.125.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.126.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.126.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.126.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.127.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.127.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.127.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.128.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.128.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.128.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.129.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.129.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.129.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.130.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.130.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.130.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.131.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.131.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.131.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.132.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.132.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.132.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.133.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.133.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.133.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.134.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.134.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.134.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.135.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.135.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.135.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.136.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.136.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.136.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.137.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.137.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.137.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.138.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.138.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.138.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.139.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.139.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.139.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.140.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.140.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.140.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.141.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.141.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.141.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.142.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.142.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.142.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.143.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.143.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.143.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.144.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.144.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.144.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.145.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.145.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.145.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.146.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.146.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.146.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.147.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.147.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.147.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.148.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.148.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.148.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.149.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.149.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.149.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.150.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.150.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.150.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.151.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.151.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.151.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.152.gate_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.152.up_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.152.down_proj.weight": "model-00027-of-000163.safetensors", + "model.layers.12.mlp.experts.153.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.153.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.153.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.154.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.154.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.154.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.155.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.155.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.155.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.156.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.156.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.156.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.157.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.157.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.157.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.158.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.158.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.158.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.159.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.159.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.159.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.160.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.160.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.160.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.161.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.161.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.161.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.162.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.162.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.162.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.163.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.163.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.163.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.164.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.164.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.164.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.165.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.165.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.165.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.166.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.166.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.166.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.167.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.167.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.167.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.168.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.168.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.168.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.169.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.169.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.169.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.170.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.170.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.170.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.171.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.171.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.171.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.172.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.172.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.172.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.173.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.173.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.173.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.174.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.174.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.174.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.175.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.175.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.175.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.176.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.176.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.176.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.177.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.177.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.177.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.178.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.178.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.178.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.179.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.179.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.179.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.180.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.180.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.180.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.181.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.181.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.181.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.182.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.182.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.182.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.183.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.183.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.183.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.184.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.184.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.184.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.185.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.185.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.185.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.186.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.186.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.186.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.187.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.187.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.187.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.188.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.188.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.188.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.189.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.189.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.189.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.190.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.190.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.190.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.191.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.191.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.191.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.192.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.192.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.192.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.193.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.193.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.193.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.194.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.194.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.194.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.195.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.195.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.195.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.196.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.196.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.196.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.197.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.197.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.197.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.198.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.198.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.198.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.199.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.199.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.199.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.200.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.200.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.200.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.201.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.201.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.201.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.202.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.202.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.202.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.203.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.203.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.203.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.204.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.204.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.204.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.205.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.205.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.205.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.206.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.206.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.206.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.207.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.207.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.207.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.208.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.208.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.208.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.209.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.209.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.209.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.210.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.210.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.210.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.211.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.211.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.211.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.212.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.212.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.212.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.213.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.213.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.213.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.214.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.214.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.214.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.215.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.215.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.215.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.216.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.216.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.216.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.217.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.217.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.217.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.218.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.218.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.218.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.219.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.219.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.219.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.220.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.220.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.220.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.221.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.221.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.221.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.222.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.222.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.222.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.223.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.223.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.223.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.224.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.224.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.224.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.225.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.225.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.225.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.226.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.226.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.226.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.227.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.227.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.227.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.228.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.228.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.228.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.229.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.229.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.229.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.230.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.230.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.230.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.231.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.231.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.231.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.232.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.232.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.232.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.233.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.233.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.233.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.234.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.234.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.234.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.235.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.235.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.235.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.236.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.236.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.236.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.237.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.237.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.237.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.238.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.238.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.238.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.239.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.239.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.239.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.240.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.240.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.240.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.241.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.241.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.241.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.242.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.242.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.242.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.243.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.243.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.243.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.244.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.244.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.244.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.245.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.245.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.245.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.246.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.246.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.246.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.247.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.247.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.247.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.248.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.248.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.248.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.249.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.249.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.249.down_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.250.gate_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.250.up_proj.weight": "model-00028-of-000163.safetensors", + "model.layers.12.mlp.experts.250.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.12.mlp.experts.251.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.12.mlp.experts.251.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.12.mlp.experts.251.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.12.mlp.experts.252.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.12.mlp.experts.252.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.12.mlp.experts.252.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.12.mlp.experts.253.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.12.mlp.experts.253.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.12.mlp.experts.253.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.12.mlp.experts.254.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.12.mlp.experts.254.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.12.mlp.experts.254.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.12.mlp.experts.255.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.12.mlp.experts.255.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.12.mlp.experts.255.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.12.input_layernorm.weight": "model-00029-of-000163.safetensors", + "model.layers.12.post_attention_layernorm.weight": "model-00029-of-000163.safetensors", + "model.layers.13.self_attn.q_a_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.self_attn.q_a_layernorm.weight": "model-00029-of-000163.safetensors", + "model.layers.13.self_attn.q_b_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.self_attn.kv_a_proj_with_mqa.weight": "model-00029-of-000163.safetensors", + "model.layers.13.self_attn.kv_a_layernorm.weight": "model-00029-of-000163.safetensors", + "model.layers.13.self_attn.kv_b_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.self_attn.o_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.gate.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.gate.e_score_correction_bias": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.shared_experts.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.shared_experts.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.shared_experts.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.0.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.0.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.0.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.1.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.1.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.1.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.2.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.2.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.2.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.3.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.3.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.3.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.4.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.4.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.4.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.5.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.5.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.5.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.6.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.6.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.6.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.7.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.7.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.7.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.8.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.8.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.8.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.9.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.9.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.9.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.10.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.10.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.10.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.11.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.11.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.11.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.12.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.12.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.12.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.13.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.13.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.13.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.14.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.14.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.14.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.15.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.15.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.15.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.16.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.16.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.16.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.17.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.17.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.17.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.18.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.18.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.18.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.19.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.19.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.19.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.20.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.20.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.20.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.21.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.21.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.21.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.22.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.22.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.22.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.23.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.23.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.23.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.24.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.24.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.24.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.25.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.25.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.25.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.26.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.26.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.26.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.27.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.27.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.27.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.28.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.28.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.28.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.29.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.29.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.29.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.30.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.30.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.30.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.31.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.31.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.31.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.32.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.32.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.32.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.33.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.33.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.33.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.34.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.34.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.34.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.35.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.35.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.35.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.36.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.36.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.36.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.37.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.37.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.37.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.38.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.38.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.38.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.39.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.39.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.39.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.40.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.40.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.40.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.41.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.41.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.41.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.42.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.42.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.42.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.43.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.43.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.43.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.44.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.44.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.44.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.45.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.45.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.45.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.46.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.46.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.46.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.47.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.47.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.47.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.48.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.48.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.48.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.49.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.49.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.49.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.50.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.50.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.50.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.51.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.51.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.51.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.52.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.52.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.52.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.53.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.53.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.53.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.54.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.54.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.54.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.55.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.55.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.55.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.56.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.56.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.56.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.57.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.57.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.57.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.58.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.58.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.58.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.59.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.59.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.59.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.60.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.60.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.60.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.61.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.61.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.61.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.62.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.62.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.62.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.63.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.63.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.63.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.64.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.64.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.64.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.65.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.65.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.65.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.66.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.66.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.66.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.67.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.67.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.67.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.68.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.68.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.68.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.69.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.69.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.69.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.70.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.70.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.70.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.71.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.71.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.71.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.72.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.72.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.72.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.73.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.73.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.73.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.74.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.74.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.74.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.75.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.75.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.75.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.76.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.76.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.76.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.77.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.77.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.77.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.78.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.78.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.78.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.79.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.79.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.79.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.80.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.80.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.80.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.81.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.81.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.81.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.82.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.82.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.82.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.83.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.83.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.83.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.84.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.84.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.84.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.85.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.85.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.85.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.86.gate_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.86.up_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.86.down_proj.weight": "model-00029-of-000163.safetensors", + "model.layers.13.mlp.experts.87.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.87.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.87.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.88.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.88.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.88.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.89.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.89.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.89.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.90.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.90.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.90.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.91.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.91.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.91.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.92.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.92.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.92.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.93.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.93.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.93.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.94.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.94.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.94.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.95.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.95.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.95.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.96.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.96.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.96.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.97.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.97.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.97.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.98.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.98.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.98.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.99.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.99.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.99.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.100.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.100.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.100.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.101.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.101.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.101.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.102.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.102.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.102.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.103.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.103.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.103.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.104.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.104.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.104.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.105.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.105.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.105.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.106.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.106.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.106.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.107.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.107.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.107.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.108.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.108.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.108.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.109.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.109.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.109.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.110.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.110.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.110.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.111.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.111.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.111.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.112.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.112.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.112.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.113.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.113.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.113.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.114.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.114.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.114.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.115.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.115.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.115.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.116.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.116.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.116.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.117.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.117.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.117.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.118.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.118.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.118.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.119.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.119.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.119.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.120.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.120.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.120.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.121.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.121.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.121.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.122.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.122.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.122.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.123.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.123.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.123.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.124.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.124.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.124.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.125.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.125.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.125.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.126.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.126.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.126.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.127.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.127.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.127.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.128.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.128.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.128.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.129.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.129.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.129.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.130.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.130.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.130.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.131.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.131.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.131.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.132.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.132.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.132.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.133.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.133.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.133.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.134.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.134.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.134.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.135.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.135.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.135.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.136.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.136.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.136.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.137.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.137.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.137.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.138.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.138.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.138.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.139.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.139.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.139.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.140.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.140.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.140.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.141.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.141.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.141.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.142.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.142.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.142.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.143.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.143.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.143.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.144.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.144.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.144.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.145.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.145.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.145.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.146.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.146.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.146.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.147.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.147.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.147.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.148.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.148.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.148.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.149.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.149.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.149.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.150.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.150.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.150.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.151.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.151.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.151.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.152.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.152.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.152.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.153.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.153.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.153.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.154.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.154.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.154.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.155.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.155.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.155.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.156.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.156.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.156.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.157.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.157.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.157.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.158.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.158.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.158.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.159.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.159.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.159.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.160.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.160.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.160.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.161.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.161.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.161.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.162.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.162.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.162.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.163.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.163.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.163.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.164.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.164.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.164.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.165.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.165.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.165.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.166.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.166.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.166.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.167.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.167.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.167.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.168.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.168.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.168.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.169.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.169.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.169.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.170.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.170.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.170.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.171.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.171.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.171.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.172.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.172.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.172.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.173.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.173.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.173.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.174.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.174.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.174.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.175.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.175.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.175.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.176.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.176.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.176.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.177.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.177.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.177.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.178.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.178.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.178.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.179.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.179.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.179.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.180.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.180.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.180.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.181.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.181.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.181.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.182.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.182.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.182.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.183.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.183.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.183.down_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.184.gate_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.184.up_proj.weight": "model-00030-of-000163.safetensors", + "model.layers.13.mlp.experts.184.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.185.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.185.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.185.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.186.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.186.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.186.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.187.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.187.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.187.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.188.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.188.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.188.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.189.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.189.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.189.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.190.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.190.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.190.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.191.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.191.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.191.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.192.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.192.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.192.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.193.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.193.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.193.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.194.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.194.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.194.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.195.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.195.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.195.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.196.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.196.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.196.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.197.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.197.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.197.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.198.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.198.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.198.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.199.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.199.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.199.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.200.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.200.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.200.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.201.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.201.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.201.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.202.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.202.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.202.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.203.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.203.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.203.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.204.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.204.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.204.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.205.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.205.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.205.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.206.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.206.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.206.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.207.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.207.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.207.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.208.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.208.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.208.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.209.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.209.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.209.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.210.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.210.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.210.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.211.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.211.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.211.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.212.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.212.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.212.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.213.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.213.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.213.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.214.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.214.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.214.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.215.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.215.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.215.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.216.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.216.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.216.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.217.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.217.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.217.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.218.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.218.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.218.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.219.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.219.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.219.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.220.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.220.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.220.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.221.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.221.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.221.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.222.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.222.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.222.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.223.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.223.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.223.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.224.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.224.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.224.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.225.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.225.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.225.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.226.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.226.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.226.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.227.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.227.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.227.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.228.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.228.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.228.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.229.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.229.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.229.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.230.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.230.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.230.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.231.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.231.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.231.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.232.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.232.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.232.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.233.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.233.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.233.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.234.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.234.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.234.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.235.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.235.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.235.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.236.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.236.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.236.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.237.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.237.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.237.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.238.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.238.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.238.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.239.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.239.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.239.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.240.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.240.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.240.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.241.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.241.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.241.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.242.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.242.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.242.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.243.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.243.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.243.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.244.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.244.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.244.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.245.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.245.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.245.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.246.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.246.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.246.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.247.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.247.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.247.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.248.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.248.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.248.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.249.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.249.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.249.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.250.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.250.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.250.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.251.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.251.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.251.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.252.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.252.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.252.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.253.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.253.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.253.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.254.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.254.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.254.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.255.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.255.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.mlp.experts.255.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.13.input_layernorm.weight": "model-00031-of-000163.safetensors", + "model.layers.13.post_attention_layernorm.weight": "model-00031-of-000163.safetensors", + "model.layers.14.self_attn.q_a_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.self_attn.q_a_layernorm.weight": "model-00031-of-000163.safetensors", + "model.layers.14.self_attn.q_b_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.self_attn.kv_a_proj_with_mqa.weight": "model-00031-of-000163.safetensors", + "model.layers.14.self_attn.kv_a_layernorm.weight": "model-00031-of-000163.safetensors", + "model.layers.14.self_attn.kv_b_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.self_attn.o_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.gate.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.gate.e_score_correction_bias": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.shared_experts.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.shared_experts.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.shared_experts.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.0.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.0.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.0.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.1.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.1.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.1.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.2.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.2.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.2.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.3.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.3.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.3.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.4.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.4.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.4.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.5.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.5.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.5.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.6.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.6.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.6.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.7.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.7.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.7.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.8.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.8.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.8.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.9.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.9.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.9.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.10.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.10.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.10.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.11.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.11.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.11.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.12.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.12.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.12.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.13.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.13.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.13.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.14.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.14.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.14.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.15.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.15.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.15.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.16.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.16.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.16.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.17.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.17.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.17.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.18.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.18.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.18.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.19.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.19.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.19.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.20.gate_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.20.up_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.20.down_proj.weight": "model-00031-of-000163.safetensors", + "model.layers.14.mlp.experts.21.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.21.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.21.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.22.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.22.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.22.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.23.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.23.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.23.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.24.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.24.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.24.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.25.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.25.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.25.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.26.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.26.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.26.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.27.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.27.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.27.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.28.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.28.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.28.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.29.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.29.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.29.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.30.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.30.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.30.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.31.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.31.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.31.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.32.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.32.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.32.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.33.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.33.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.33.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.34.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.34.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.34.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.35.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.35.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.35.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.36.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.36.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.36.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.37.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.37.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.37.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.38.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.38.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.38.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.39.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.39.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.39.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.40.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.40.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.40.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.41.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.41.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.41.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.42.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.42.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.42.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.43.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.43.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.43.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.44.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.44.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.44.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.45.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.45.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.45.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.46.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.46.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.46.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.47.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.47.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.47.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.48.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.48.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.48.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.49.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.49.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.49.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.50.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.50.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.50.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.51.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.51.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.51.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.52.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.52.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.52.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.53.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.53.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.53.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.54.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.54.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.54.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.55.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.55.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.55.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.56.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.56.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.56.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.57.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.57.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.57.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.58.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.58.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.58.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.59.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.59.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.59.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.60.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.60.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.60.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.61.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.61.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.61.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.62.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.62.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.62.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.63.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.63.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.63.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.64.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.64.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.64.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.65.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.65.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.65.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.66.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.66.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.66.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.67.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.67.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.67.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.68.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.68.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.68.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.69.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.69.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.69.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.70.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.70.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.70.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.71.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.71.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.71.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.72.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.72.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.72.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.73.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.73.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.73.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.74.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.74.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.74.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.75.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.75.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.75.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.76.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.76.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.76.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.77.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.77.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.77.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.78.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.78.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.78.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.79.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.79.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.79.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.80.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.80.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.80.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.81.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.81.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.81.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.82.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.82.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.82.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.83.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.83.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.83.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.84.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.84.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.84.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.85.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.85.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.85.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.86.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.86.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.86.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.87.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.87.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.87.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.88.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.88.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.88.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.89.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.89.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.89.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.90.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.90.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.90.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.91.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.91.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.91.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.92.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.92.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.92.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.93.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.93.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.93.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.94.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.94.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.94.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.95.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.95.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.95.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.96.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.96.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.96.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.97.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.97.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.97.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.98.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.98.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.98.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.99.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.99.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.99.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.100.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.100.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.100.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.101.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.101.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.101.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.102.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.102.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.102.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.103.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.103.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.103.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.104.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.104.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.104.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.105.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.105.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.105.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.106.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.106.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.106.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.107.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.107.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.107.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.108.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.108.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.108.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.109.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.109.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.109.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.110.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.110.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.110.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.111.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.111.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.111.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.112.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.112.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.112.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.113.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.113.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.113.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.114.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.114.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.114.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.115.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.115.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.115.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.116.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.116.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.116.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.117.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.117.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.117.down_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.118.gate_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.118.up_proj.weight": "model-00032-of-000163.safetensors", + "model.layers.14.mlp.experts.118.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.119.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.119.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.119.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.120.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.120.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.120.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.121.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.121.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.121.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.122.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.122.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.122.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.123.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.123.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.123.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.124.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.124.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.124.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.125.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.125.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.125.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.126.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.126.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.126.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.127.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.127.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.127.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.128.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.128.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.128.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.129.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.129.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.129.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.130.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.130.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.130.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.131.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.131.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.131.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.132.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.132.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.132.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.133.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.133.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.133.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.134.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.134.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.134.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.135.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.135.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.135.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.136.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.136.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.136.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.137.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.137.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.137.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.138.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.138.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.138.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.139.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.139.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.139.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.140.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.140.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.140.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.141.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.141.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.141.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.142.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.142.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.142.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.143.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.143.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.143.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.144.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.144.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.144.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.145.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.145.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.145.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.146.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.146.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.146.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.147.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.147.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.147.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.148.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.148.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.148.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.149.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.149.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.149.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.150.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.150.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.150.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.151.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.151.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.151.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.152.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.152.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.152.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.153.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.153.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.153.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.154.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.154.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.154.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.155.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.155.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.155.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.156.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.156.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.156.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.157.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.157.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.157.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.158.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.158.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.158.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.159.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.159.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.159.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.160.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.160.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.160.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.161.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.161.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.161.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.162.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.162.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.162.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.163.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.163.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.163.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.164.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.164.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.164.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.165.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.165.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.165.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.166.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.166.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.166.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.167.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.167.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.167.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.168.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.168.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.168.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.169.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.169.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.169.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.170.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.170.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.170.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.171.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.171.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.171.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.172.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.172.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.172.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.173.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.173.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.173.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.174.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.174.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.174.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.175.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.175.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.175.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.176.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.176.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.176.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.177.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.177.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.177.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.178.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.178.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.178.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.179.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.179.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.179.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.180.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.180.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.180.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.181.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.181.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.181.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.182.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.182.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.182.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.183.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.183.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.183.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.184.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.184.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.184.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.185.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.185.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.185.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.186.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.186.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.186.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.187.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.187.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.187.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.188.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.188.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.188.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.189.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.189.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.189.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.190.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.190.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.190.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.191.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.191.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.191.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.192.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.192.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.192.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.193.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.193.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.193.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.194.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.194.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.194.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.195.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.195.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.195.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.196.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.196.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.196.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.197.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.197.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.197.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.198.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.198.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.198.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.199.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.199.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.199.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.200.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.200.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.200.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.201.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.201.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.201.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.202.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.202.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.202.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.203.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.203.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.203.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.204.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.204.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.204.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.205.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.205.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.205.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.206.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.206.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.206.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.207.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.207.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.207.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.208.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.208.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.208.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.209.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.209.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.209.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.210.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.210.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.210.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.211.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.211.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.211.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.212.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.212.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.212.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.213.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.213.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.213.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.214.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.214.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.214.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.215.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.215.up_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.215.down_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.216.gate_proj.weight": "model-00033-of-000163.safetensors", + "model.layers.14.mlp.experts.216.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.216.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.217.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.217.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.217.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.218.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.218.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.218.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.219.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.219.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.219.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.220.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.220.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.220.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.221.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.221.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.221.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.222.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.222.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.222.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.223.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.223.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.223.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.224.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.224.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.224.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.225.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.225.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.225.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.226.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.226.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.226.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.227.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.227.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.227.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.228.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.228.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.228.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.229.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.229.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.229.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.230.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.230.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.230.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.231.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.231.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.231.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.232.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.232.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.232.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.233.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.233.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.233.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.234.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.234.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.234.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.235.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.235.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.235.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.236.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.236.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.236.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.237.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.237.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.237.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.238.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.238.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.238.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.239.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.239.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.239.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.240.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.240.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.240.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.241.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.241.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.241.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.242.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.242.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.242.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.243.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.243.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.243.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.244.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.244.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.244.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.245.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.245.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.245.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.246.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.246.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.246.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.247.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.247.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.247.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.248.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.248.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.248.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.249.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.249.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.249.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.250.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.250.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.250.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.251.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.251.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.251.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.252.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.252.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.252.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.253.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.253.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.253.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.254.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.254.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.254.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.255.gate_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.255.up_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.mlp.experts.255.down_proj.weight": "model-00034-of-000163.safetensors", + "model.layers.14.input_layernorm.weight": "model-00034-of-000163.safetensors", + "model.layers.14.post_attention_layernorm.weight": "model-00034-of-000163.safetensors", + "model.layers.15.self_attn.q_a_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.self_attn.q_a_layernorm.weight": "model-00035-of-000163.safetensors", + "model.layers.15.self_attn.q_b_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.self_attn.kv_a_proj_with_mqa.weight": "model-00035-of-000163.safetensors", + "model.layers.15.self_attn.kv_a_layernorm.weight": "model-00035-of-000163.safetensors", + "model.layers.15.self_attn.kv_b_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.self_attn.o_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.gate.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.gate.e_score_correction_bias": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.shared_experts.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.shared_experts.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.shared_experts.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.0.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.0.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.0.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.1.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.1.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.1.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.2.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.2.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.2.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.3.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.3.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.3.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.4.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.4.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.4.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.5.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.5.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.5.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.6.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.6.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.6.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.7.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.7.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.7.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.8.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.8.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.8.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.9.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.9.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.9.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.10.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.10.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.10.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.11.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.11.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.11.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.12.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.12.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.12.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.13.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.13.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.13.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.14.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.14.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.14.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.15.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.15.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.15.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.16.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.16.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.16.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.17.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.17.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.17.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.18.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.18.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.18.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.19.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.19.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.19.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.20.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.20.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.20.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.21.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.21.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.21.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.22.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.22.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.22.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.23.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.23.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.23.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.24.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.24.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.24.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.25.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.25.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.25.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.26.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.26.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.26.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.27.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.27.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.27.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.28.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.28.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.28.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.29.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.29.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.29.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.30.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.30.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.30.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.31.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.31.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.31.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.32.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.32.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.32.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.33.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.33.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.33.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.34.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.34.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.34.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.35.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.35.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.35.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.36.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.36.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.36.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.37.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.37.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.37.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.38.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.38.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.38.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.39.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.39.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.39.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.40.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.40.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.40.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.41.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.41.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.41.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.42.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.42.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.42.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.43.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.43.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.43.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.44.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.44.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.44.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.45.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.45.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.45.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.46.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.46.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.46.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.47.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.47.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.47.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.48.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.48.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.48.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.49.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.49.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.49.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.50.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.50.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.50.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.51.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.51.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.51.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.52.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.52.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.52.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.53.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.53.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.53.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.54.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.54.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.54.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.55.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.55.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.55.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.56.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.56.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.56.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.57.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.57.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.57.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.58.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.58.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.58.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.59.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.59.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.59.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.60.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.60.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.60.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.61.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.61.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.61.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.62.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.62.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.62.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.63.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.63.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.63.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.64.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.64.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.64.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.65.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.65.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.65.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.66.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.66.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.66.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.67.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.67.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.67.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.68.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.68.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.68.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.69.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.69.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.69.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.70.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.70.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.70.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.71.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.71.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.71.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.72.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.72.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.72.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.73.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.73.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.73.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.74.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.74.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.74.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.75.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.75.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.75.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.76.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.76.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.76.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.77.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.77.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.77.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.78.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.78.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.78.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.79.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.79.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.79.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.80.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.80.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.80.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.81.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.81.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.81.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.82.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.82.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.82.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.83.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.83.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.83.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.84.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.84.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.84.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.85.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.85.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.85.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.86.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.86.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.86.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.87.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.87.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.87.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.88.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.88.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.88.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.89.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.89.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.89.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.90.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.90.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.90.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.91.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.91.up_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.91.down_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.92.gate_proj.weight": "model-00035-of-000163.safetensors", + "model.layers.15.mlp.experts.92.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.92.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.93.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.93.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.93.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.94.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.94.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.94.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.95.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.95.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.95.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.96.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.96.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.96.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.97.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.97.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.97.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.98.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.98.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.98.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.99.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.99.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.99.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.100.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.100.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.100.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.101.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.101.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.101.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.102.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.102.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.102.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.103.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.103.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.103.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.104.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.104.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.104.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.105.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.105.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.105.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.106.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.106.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.106.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.107.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.107.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.107.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.108.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.108.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.108.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.109.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.109.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.109.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.110.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.110.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.110.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.111.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.111.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.111.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.112.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.112.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.112.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.113.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.113.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.113.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.114.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.114.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.114.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.115.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.115.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.115.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.116.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.116.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.116.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.117.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.117.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.117.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.118.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.118.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.118.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.119.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.119.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.119.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.120.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.120.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.120.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.121.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.121.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.121.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.122.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.122.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.122.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.123.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.123.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.123.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.124.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.124.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.124.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.125.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.125.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.125.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.126.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.126.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.126.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.127.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.127.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.127.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.128.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.128.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.128.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.129.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.129.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.129.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.130.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.130.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.130.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.131.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.131.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.131.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.132.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.132.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.132.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.133.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.133.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.133.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.134.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.134.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.134.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.135.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.135.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.135.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.136.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.136.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.136.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.137.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.137.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.137.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.138.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.138.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.138.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.139.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.139.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.139.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.140.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.140.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.140.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.141.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.141.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.141.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.142.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.142.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.142.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.143.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.143.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.143.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.144.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.144.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.144.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.145.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.145.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.145.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.146.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.146.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.146.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.147.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.147.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.147.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.148.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.148.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.148.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.149.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.149.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.149.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.150.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.150.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.150.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.151.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.151.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.151.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.152.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.152.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.152.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.153.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.153.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.153.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.154.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.154.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.154.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.155.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.155.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.155.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.156.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.156.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.156.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.157.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.157.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.157.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.158.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.158.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.158.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.159.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.159.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.159.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.160.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.160.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.160.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.161.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.161.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.161.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.162.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.162.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.162.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.163.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.163.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.163.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.164.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.164.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.164.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.165.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.165.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.165.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.166.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.166.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.166.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.167.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.167.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.167.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.168.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.168.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.168.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.169.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.169.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.169.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.170.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.170.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.170.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.171.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.171.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.171.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.172.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.172.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.172.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.173.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.173.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.173.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.174.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.174.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.174.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.175.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.175.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.175.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.176.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.176.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.176.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.177.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.177.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.177.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.178.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.178.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.178.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.179.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.179.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.179.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.180.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.180.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.180.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.181.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.181.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.181.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.182.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.182.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.182.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.183.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.183.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.183.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.184.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.184.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.184.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.185.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.185.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.185.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.186.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.186.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.186.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.187.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.187.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.187.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.188.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.188.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.188.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.189.gate_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.189.up_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.189.down_proj.weight": "model-00036-of-000163.safetensors", + "model.layers.15.mlp.experts.190.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.190.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.190.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.191.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.191.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.191.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.192.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.192.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.192.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.193.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.193.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.193.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.194.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.194.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.194.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.195.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.195.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.195.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.196.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.196.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.196.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.197.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.197.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.197.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.198.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.198.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.198.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.199.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.199.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.199.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.200.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.200.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.200.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.201.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.201.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.201.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.202.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.202.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.202.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.203.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.203.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.203.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.204.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.204.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.204.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.205.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.205.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.205.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.206.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.206.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.206.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.207.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.207.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.207.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.208.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.208.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.208.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.209.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.209.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.209.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.210.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.210.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.210.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.211.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.211.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.211.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.212.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.212.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.212.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.213.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.213.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.213.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.214.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.214.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.214.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.215.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.215.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.215.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.216.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.216.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.216.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.217.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.217.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.217.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.218.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.218.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.218.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.219.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.219.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.219.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.220.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.220.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.220.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.221.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.221.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.221.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.222.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.222.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.222.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.223.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.223.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.223.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.224.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.224.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.224.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.225.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.225.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.225.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.226.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.226.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.226.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.227.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.227.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.227.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.228.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.228.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.228.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.229.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.229.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.229.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.230.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.230.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.230.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.231.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.231.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.231.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.232.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.232.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.232.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.233.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.233.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.233.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.234.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.234.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.234.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.235.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.235.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.235.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.236.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.236.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.236.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.237.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.237.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.237.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.238.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.238.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.238.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.239.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.239.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.239.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.240.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.240.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.240.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.241.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.241.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.241.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.242.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.242.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.242.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.243.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.243.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.243.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.244.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.244.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.244.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.245.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.245.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.245.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.246.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.246.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.246.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.247.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.247.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.247.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.248.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.248.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.248.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.249.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.249.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.249.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.250.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.250.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.250.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.251.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.251.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.251.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.252.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.252.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.252.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.253.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.253.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.253.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.254.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.254.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.254.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.255.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.255.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.mlp.experts.255.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.15.input_layernorm.weight": "model-00037-of-000163.safetensors", + "model.layers.15.post_attention_layernorm.weight": "model-00037-of-000163.safetensors", + "model.layers.16.self_attn.q_a_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.self_attn.q_a_layernorm.weight": "model-00037-of-000163.safetensors", + "model.layers.16.self_attn.q_b_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.self_attn.kv_a_proj_with_mqa.weight": "model-00037-of-000163.safetensors", + "model.layers.16.self_attn.kv_a_layernorm.weight": "model-00037-of-000163.safetensors", + "model.layers.16.self_attn.kv_b_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.self_attn.o_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.gate.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.gate.e_score_correction_bias": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.shared_experts.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.shared_experts.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.shared_experts.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.0.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.0.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.0.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.1.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.1.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.1.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.2.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.2.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.2.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.3.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.3.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.3.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.4.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.4.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.4.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.5.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.5.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.5.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.6.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.6.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.6.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.7.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.7.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.7.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.8.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.8.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.8.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.9.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.9.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.9.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.10.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.10.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.10.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.11.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.11.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.11.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.12.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.12.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.12.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.13.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.13.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.13.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.14.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.14.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.14.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.15.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.15.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.15.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.16.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.16.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.16.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.17.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.17.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.17.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.18.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.18.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.18.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.19.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.19.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.19.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.20.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.20.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.20.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.21.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.21.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.21.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.22.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.22.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.22.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.23.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.23.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.23.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.24.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.24.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.24.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.25.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.25.up_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.25.down_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.26.gate_proj.weight": "model-00037-of-000163.safetensors", + "model.layers.16.mlp.experts.26.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.26.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.27.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.27.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.27.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.28.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.28.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.28.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.29.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.29.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.29.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.30.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.30.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.30.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.31.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.31.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.31.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.32.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.32.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.32.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.33.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.33.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.33.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.34.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.34.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.34.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.35.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.35.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.35.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.36.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.36.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.36.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.37.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.37.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.37.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.38.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.38.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.38.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.39.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.39.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.39.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.40.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.40.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.40.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.41.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.41.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.41.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.42.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.42.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.42.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.43.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.43.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.43.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.44.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.44.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.44.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.45.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.45.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.45.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.46.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.46.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.46.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.47.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.47.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.47.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.48.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.48.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.48.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.49.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.49.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.49.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.50.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.50.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.50.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.51.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.51.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.51.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.52.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.52.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.52.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.53.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.53.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.53.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.54.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.54.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.54.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.55.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.55.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.55.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.56.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.56.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.56.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.57.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.57.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.57.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.58.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.58.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.58.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.59.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.59.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.59.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.60.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.60.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.60.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.61.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.61.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.61.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.62.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.62.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.62.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.63.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.63.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.63.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.64.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.64.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.64.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.65.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.65.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.65.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.66.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.66.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.66.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.67.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.67.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.67.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.68.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.68.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.68.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.69.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.69.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.69.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.70.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.70.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.70.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.71.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.71.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.71.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.72.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.72.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.72.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.73.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.73.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.73.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.74.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.74.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.74.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.75.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.75.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.75.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.76.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.76.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.76.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.77.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.77.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.77.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.78.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.78.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.78.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.79.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.79.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.79.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.80.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.80.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.80.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.81.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.81.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.81.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.82.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.82.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.82.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.83.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.83.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.83.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.84.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.84.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.84.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.85.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.85.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.85.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.86.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.86.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.86.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.87.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.87.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.87.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.88.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.88.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.88.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.89.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.89.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.89.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.90.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.90.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.90.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.91.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.91.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.91.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.92.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.92.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.92.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.93.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.93.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.93.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.94.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.94.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.94.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.95.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.95.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.95.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.96.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.96.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.96.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.97.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.97.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.97.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.98.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.98.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.98.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.99.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.99.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.99.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.100.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.100.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.100.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.101.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.101.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.101.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.102.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.102.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.102.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.103.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.103.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.103.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.104.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.104.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.104.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.105.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.105.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.105.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.106.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.106.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.106.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.107.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.107.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.107.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.108.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.108.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.108.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.109.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.109.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.109.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.110.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.110.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.110.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.111.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.111.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.111.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.112.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.112.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.112.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.113.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.113.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.113.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.114.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.114.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.114.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.115.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.115.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.115.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.116.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.116.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.116.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.117.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.117.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.117.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.118.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.118.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.118.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.119.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.119.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.119.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.120.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.120.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.120.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.121.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.121.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.121.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.122.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.122.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.122.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.123.gate_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.123.up_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.123.down_proj.weight": "model-00038-of-000163.safetensors", + "model.layers.16.mlp.experts.124.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.124.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.124.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.125.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.125.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.125.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.126.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.126.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.126.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.127.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.127.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.127.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.128.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.128.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.128.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.129.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.129.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.129.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.130.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.130.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.130.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.131.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.131.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.131.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.132.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.132.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.132.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.133.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.133.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.133.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.134.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.134.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.134.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.135.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.135.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.135.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.136.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.136.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.136.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.137.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.137.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.137.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.138.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.138.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.138.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.139.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.139.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.139.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.140.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.140.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.140.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.141.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.141.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.141.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.142.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.142.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.142.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.143.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.143.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.143.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.144.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.144.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.144.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.145.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.145.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.145.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.146.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.146.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.146.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.147.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.147.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.147.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.148.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.148.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.148.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.149.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.149.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.149.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.150.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.150.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.150.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.151.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.151.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.151.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.152.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.152.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.152.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.153.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.153.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.153.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.154.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.154.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.154.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.155.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.155.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.155.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.156.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.156.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.156.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.157.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.157.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.157.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.158.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.158.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.158.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.159.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.159.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.159.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.160.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.160.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.160.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.161.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.161.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.161.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.162.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.162.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.162.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.163.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.163.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.163.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.164.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.164.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.164.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.165.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.165.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.165.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.166.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.166.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.166.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.167.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.167.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.167.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.168.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.168.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.168.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.169.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.169.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.169.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.170.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.170.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.170.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.171.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.171.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.171.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.172.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.172.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.172.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.173.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.173.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.173.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.174.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.174.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.174.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.175.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.175.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.175.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.176.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.176.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.176.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.177.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.177.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.177.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.178.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.178.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.178.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.179.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.179.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.179.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.180.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.180.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.180.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.181.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.181.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.181.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.182.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.182.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.182.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.183.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.183.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.183.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.184.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.184.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.184.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.185.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.185.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.185.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.186.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.186.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.186.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.187.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.187.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.187.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.188.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.188.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.188.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.189.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.189.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.189.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.190.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.190.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.190.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.191.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.191.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.191.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.192.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.192.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.192.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.193.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.193.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.193.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.194.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.194.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.194.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.195.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.195.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.195.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.196.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.196.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.196.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.197.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.197.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.197.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.198.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.198.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.198.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.199.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.199.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.199.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.200.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.200.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.200.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.201.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.201.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.201.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.202.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.202.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.202.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.203.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.203.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.203.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.204.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.204.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.204.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.205.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.205.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.205.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.206.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.206.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.206.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.207.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.207.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.207.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.208.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.208.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.208.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.209.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.209.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.209.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.210.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.210.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.210.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.211.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.211.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.211.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.212.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.212.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.212.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.213.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.213.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.213.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.214.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.214.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.214.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.215.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.215.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.215.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.216.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.216.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.216.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.217.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.217.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.217.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.218.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.218.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.218.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.219.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.219.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.219.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.220.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.220.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.220.down_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.221.gate_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.221.up_proj.weight": "model-00039-of-000163.safetensors", + "model.layers.16.mlp.experts.221.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.222.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.222.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.222.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.223.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.223.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.223.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.224.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.224.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.224.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.225.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.225.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.225.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.226.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.226.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.226.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.227.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.227.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.227.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.228.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.228.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.228.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.229.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.229.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.229.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.230.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.230.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.230.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.231.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.231.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.231.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.232.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.232.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.232.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.233.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.233.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.233.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.234.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.234.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.234.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.235.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.235.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.235.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.236.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.236.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.236.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.237.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.237.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.237.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.238.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.238.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.238.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.239.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.239.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.239.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.240.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.240.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.240.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.241.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.241.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.241.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.242.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.242.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.242.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.243.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.243.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.243.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.244.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.244.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.244.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.245.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.245.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.245.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.246.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.246.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.246.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.247.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.247.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.247.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.248.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.248.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.248.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.249.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.249.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.249.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.250.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.250.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.250.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.251.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.251.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.251.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.252.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.252.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.252.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.253.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.253.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.253.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.254.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.254.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.254.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.255.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.255.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.mlp.experts.255.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.16.input_layernorm.weight": "model-00040-of-000163.safetensors", + "model.layers.16.post_attention_layernorm.weight": "model-00040-of-000163.safetensors", + "model.layers.17.self_attn.q_a_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.self_attn.q_a_layernorm.weight": "model-00040-of-000163.safetensors", + "model.layers.17.self_attn.q_b_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.self_attn.kv_a_proj_with_mqa.weight": "model-00040-of-000163.safetensors", + "model.layers.17.self_attn.kv_a_layernorm.weight": "model-00040-of-000163.safetensors", + "model.layers.17.self_attn.kv_b_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.self_attn.o_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.gate.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.gate.e_score_correction_bias": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.shared_experts.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.shared_experts.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.shared_experts.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.0.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.0.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.0.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.1.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.1.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.1.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.2.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.2.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.2.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.3.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.3.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.3.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.4.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.4.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.4.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.5.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.5.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.5.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.6.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.6.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.6.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.7.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.7.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.7.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.8.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.8.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.8.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.9.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.9.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.9.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.10.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.10.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.10.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.11.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.11.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.11.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.12.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.12.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.12.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.13.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.13.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.13.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.14.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.14.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.14.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.15.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.15.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.15.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.16.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.16.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.16.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.17.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.17.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.17.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.18.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.18.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.18.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.19.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.19.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.19.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.20.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.20.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.20.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.21.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.21.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.21.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.22.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.22.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.22.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.23.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.23.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.23.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.24.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.24.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.24.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.25.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.25.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.25.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.26.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.26.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.26.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.27.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.27.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.27.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.28.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.28.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.28.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.29.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.29.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.29.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.30.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.30.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.30.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.31.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.31.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.31.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.32.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.32.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.32.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.33.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.33.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.33.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.34.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.34.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.34.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.35.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.35.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.35.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.36.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.36.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.36.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.37.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.37.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.37.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.38.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.38.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.38.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.39.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.39.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.39.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.40.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.40.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.40.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.41.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.41.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.41.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.42.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.42.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.42.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.43.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.43.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.43.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.44.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.44.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.44.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.45.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.45.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.45.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.46.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.46.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.46.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.47.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.47.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.47.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.48.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.48.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.48.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.49.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.49.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.49.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.50.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.50.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.50.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.51.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.51.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.51.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.52.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.52.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.52.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.53.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.53.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.53.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.54.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.54.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.54.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.55.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.55.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.55.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.56.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.56.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.56.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.57.gate_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.57.up_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.57.down_proj.weight": "model-00040-of-000163.safetensors", + "model.layers.17.mlp.experts.58.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.58.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.58.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.59.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.59.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.59.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.60.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.60.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.60.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.61.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.61.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.61.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.62.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.62.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.62.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.63.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.63.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.63.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.64.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.64.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.64.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.65.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.65.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.65.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.66.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.66.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.66.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.67.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.67.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.67.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.68.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.68.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.68.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.69.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.69.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.69.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.70.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.70.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.70.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.71.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.71.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.71.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.72.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.72.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.72.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.73.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.73.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.73.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.74.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.74.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.74.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.75.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.75.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.75.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.76.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.76.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.76.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.77.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.77.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.77.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.78.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.78.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.78.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.79.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.79.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.79.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.80.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.80.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.80.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.81.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.81.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.81.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.82.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.82.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.82.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.83.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.83.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.83.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.84.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.84.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.84.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.85.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.85.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.85.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.86.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.86.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.86.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.87.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.87.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.87.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.88.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.88.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.88.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.89.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.89.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.89.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.90.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.90.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.90.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.91.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.91.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.91.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.92.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.92.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.92.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.93.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.93.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.93.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.94.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.94.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.94.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.95.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.95.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.95.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.96.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.96.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.96.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.97.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.97.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.97.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.98.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.98.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.98.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.99.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.99.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.99.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.100.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.100.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.100.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.101.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.101.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.101.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.102.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.102.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.102.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.103.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.103.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.103.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.104.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.104.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.104.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.105.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.105.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.105.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.106.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.106.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.106.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.107.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.107.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.107.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.108.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.108.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.108.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.109.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.109.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.109.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.110.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.110.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.110.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.111.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.111.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.111.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.112.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.112.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.112.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.113.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.113.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.113.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.114.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.114.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.114.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.115.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.115.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.115.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.116.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.116.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.116.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.117.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.117.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.117.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.118.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.118.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.118.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.119.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.119.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.119.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.120.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.120.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.120.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.121.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.121.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.121.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.122.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.122.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.122.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.123.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.123.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.123.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.124.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.124.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.124.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.125.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.125.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.125.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.126.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.126.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.126.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.127.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.127.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.127.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.128.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.128.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.128.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.129.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.129.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.129.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.130.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.130.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.130.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.131.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.131.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.131.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.132.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.132.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.132.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.133.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.133.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.133.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.134.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.134.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.134.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.135.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.135.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.135.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.136.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.136.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.136.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.137.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.137.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.137.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.138.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.138.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.138.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.139.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.139.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.139.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.140.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.140.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.140.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.141.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.141.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.141.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.142.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.142.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.142.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.143.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.143.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.143.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.144.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.144.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.144.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.145.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.145.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.145.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.146.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.146.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.146.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.147.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.147.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.147.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.148.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.148.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.148.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.149.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.149.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.149.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.150.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.150.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.150.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.151.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.151.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.151.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.152.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.152.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.152.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.153.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.153.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.153.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.154.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.154.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.154.down_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.155.gate_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.155.up_proj.weight": "model-00041-of-000163.safetensors", + "model.layers.17.mlp.experts.155.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.156.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.156.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.156.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.157.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.157.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.157.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.158.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.158.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.158.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.159.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.159.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.159.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.160.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.160.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.160.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.161.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.161.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.161.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.162.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.162.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.162.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.163.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.163.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.163.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.164.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.164.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.164.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.165.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.165.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.165.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.166.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.166.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.166.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.167.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.167.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.167.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.168.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.168.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.168.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.169.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.169.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.169.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.170.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.170.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.170.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.171.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.171.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.171.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.172.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.172.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.172.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.173.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.173.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.173.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.174.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.174.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.174.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.175.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.175.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.175.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.176.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.176.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.176.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.177.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.177.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.177.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.178.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.178.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.178.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.179.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.179.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.179.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.180.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.180.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.180.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.181.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.181.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.181.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.182.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.182.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.182.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.183.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.183.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.183.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.184.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.184.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.184.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.185.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.185.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.185.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.186.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.186.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.186.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.187.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.187.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.187.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.188.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.188.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.188.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.189.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.189.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.189.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.190.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.190.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.190.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.191.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.191.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.191.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.192.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.192.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.192.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.193.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.193.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.193.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.194.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.194.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.194.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.195.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.195.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.195.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.196.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.196.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.196.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.197.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.197.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.197.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.198.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.198.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.198.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.199.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.199.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.199.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.200.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.200.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.200.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.201.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.201.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.201.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.202.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.202.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.202.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.203.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.203.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.203.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.204.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.204.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.204.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.205.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.205.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.205.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.206.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.206.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.206.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.207.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.207.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.207.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.208.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.208.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.208.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.209.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.209.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.209.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.210.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.210.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.210.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.211.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.211.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.211.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.212.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.212.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.212.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.213.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.213.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.213.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.214.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.214.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.214.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.215.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.215.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.215.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.216.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.216.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.216.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.217.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.217.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.217.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.218.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.218.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.218.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.219.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.219.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.219.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.220.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.220.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.220.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.221.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.221.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.221.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.222.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.222.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.222.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.223.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.223.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.223.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.224.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.224.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.224.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.225.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.225.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.225.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.226.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.226.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.226.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.227.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.227.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.227.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.228.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.228.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.228.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.229.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.229.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.229.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.230.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.230.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.230.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.231.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.231.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.231.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.232.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.232.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.232.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.233.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.233.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.233.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.234.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.234.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.234.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.235.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.235.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.235.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.236.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.236.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.236.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.237.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.237.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.237.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.238.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.238.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.238.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.239.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.239.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.239.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.240.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.240.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.240.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.241.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.241.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.241.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.242.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.242.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.242.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.243.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.243.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.243.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.244.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.244.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.244.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.245.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.245.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.245.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.246.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.246.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.246.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.247.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.247.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.247.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.248.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.248.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.248.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.249.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.249.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.249.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.250.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.250.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.250.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.251.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.251.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.251.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.252.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.252.up_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.252.down_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.253.gate_proj.weight": "model-00042-of-000163.safetensors", + "model.layers.17.mlp.experts.253.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.17.mlp.experts.253.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.17.mlp.experts.254.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.17.mlp.experts.254.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.17.mlp.experts.254.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.17.mlp.experts.255.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.17.mlp.experts.255.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.17.mlp.experts.255.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.17.input_layernorm.weight": "model-00043-of-000163.safetensors", + "model.layers.17.post_attention_layernorm.weight": "model-00043-of-000163.safetensors", + "model.layers.18.self_attn.q_a_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.self_attn.q_a_layernorm.weight": "model-00043-of-000163.safetensors", + "model.layers.18.self_attn.q_b_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.self_attn.kv_a_proj_with_mqa.weight": "model-00043-of-000163.safetensors", + "model.layers.18.self_attn.kv_a_layernorm.weight": "model-00043-of-000163.safetensors", + "model.layers.18.self_attn.kv_b_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.self_attn.o_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.gate.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.gate.e_score_correction_bias": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.shared_experts.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.shared_experts.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.shared_experts.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.0.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.0.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.0.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.1.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.1.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.1.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.2.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.2.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.2.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.3.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.3.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.3.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.4.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.4.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.4.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.5.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.5.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.5.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.6.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.6.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.6.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.7.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.7.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.7.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.8.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.8.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.8.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.9.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.9.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.9.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.10.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.10.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.10.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.11.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.11.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.11.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.12.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.12.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.12.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.13.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.13.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.13.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.14.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.14.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.14.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.15.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.15.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.15.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.16.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.16.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.16.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.17.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.17.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.17.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.18.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.18.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.18.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.19.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.19.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.19.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.20.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.20.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.20.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.21.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.21.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.21.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.22.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.22.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.22.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.23.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.23.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.23.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.24.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.24.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.24.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.25.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.25.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.25.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.26.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.26.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.26.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.27.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.27.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.27.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.28.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.28.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.28.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.29.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.29.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.29.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.30.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.30.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.30.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.31.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.31.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.31.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.32.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.32.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.32.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.33.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.33.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.33.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.34.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.34.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.34.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.35.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.35.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.35.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.36.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.36.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.36.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.37.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.37.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.37.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.38.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.38.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.38.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.39.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.39.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.39.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.40.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.40.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.40.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.41.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.41.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.41.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.42.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.42.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.42.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.43.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.43.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.43.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.44.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.44.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.44.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.45.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.45.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.45.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.46.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.46.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.46.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.47.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.47.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.47.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.48.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.48.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.48.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.49.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.49.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.49.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.50.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.50.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.50.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.51.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.51.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.51.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.52.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.52.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.52.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.53.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.53.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.53.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.54.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.54.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.54.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.55.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.55.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.55.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.56.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.56.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.56.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.57.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.57.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.57.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.58.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.58.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.58.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.59.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.59.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.59.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.60.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.60.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.60.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.61.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.61.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.61.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.62.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.62.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.62.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.63.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.63.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.63.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.64.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.64.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.64.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.65.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.65.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.65.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.66.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.66.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.66.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.67.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.67.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.67.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.68.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.68.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.68.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.69.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.69.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.69.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.70.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.70.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.70.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.71.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.71.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.71.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.72.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.72.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.72.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.73.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.73.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.73.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.74.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.74.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.74.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.75.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.75.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.75.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.76.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.76.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.76.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.77.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.77.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.77.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.78.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.78.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.78.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.79.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.79.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.79.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.80.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.80.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.80.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.81.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.81.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.81.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.82.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.82.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.82.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.83.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.83.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.83.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.84.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.84.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.84.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.85.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.85.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.85.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.86.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.86.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.86.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.87.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.87.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.87.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.88.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.88.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.88.down_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.89.gate_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.89.up_proj.weight": "model-00043-of-000163.safetensors", + "model.layers.18.mlp.experts.89.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.90.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.90.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.90.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.91.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.91.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.91.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.92.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.92.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.92.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.93.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.93.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.93.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.94.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.94.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.94.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.95.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.95.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.95.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.96.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.96.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.96.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.97.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.97.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.97.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.98.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.98.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.98.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.99.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.99.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.99.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.100.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.100.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.100.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.101.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.101.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.101.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.102.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.102.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.102.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.103.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.103.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.103.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.104.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.104.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.104.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.105.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.105.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.105.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.106.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.106.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.106.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.107.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.107.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.107.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.108.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.108.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.108.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.109.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.109.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.109.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.110.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.110.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.110.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.111.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.111.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.111.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.112.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.112.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.112.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.113.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.113.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.113.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.114.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.114.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.114.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.115.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.115.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.115.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.116.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.116.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.116.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.117.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.117.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.117.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.118.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.118.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.118.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.119.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.119.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.119.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.120.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.120.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.120.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.121.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.121.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.121.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.122.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.122.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.122.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.123.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.123.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.123.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.124.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.124.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.124.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.125.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.125.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.125.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.126.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.126.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.126.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.127.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.127.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.127.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.128.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.128.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.128.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.129.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.129.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.129.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.130.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.130.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.130.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.131.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.131.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.131.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.132.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.132.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.132.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.133.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.133.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.133.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.134.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.134.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.134.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.135.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.135.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.135.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.136.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.136.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.136.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.137.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.137.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.137.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.138.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.138.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.138.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.139.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.139.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.139.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.140.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.140.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.140.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.141.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.141.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.141.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.142.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.142.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.142.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.143.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.143.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.143.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.144.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.144.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.144.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.145.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.145.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.145.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.146.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.146.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.146.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.147.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.147.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.147.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.148.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.148.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.148.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.149.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.149.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.149.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.150.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.150.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.150.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.151.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.151.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.151.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.152.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.152.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.152.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.153.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.153.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.153.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.154.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.154.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.154.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.155.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.155.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.155.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.156.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.156.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.156.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.157.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.157.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.157.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.158.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.158.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.158.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.159.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.159.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.159.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.160.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.160.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.160.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.161.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.161.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.161.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.162.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.162.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.162.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.163.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.163.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.163.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.164.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.164.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.164.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.165.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.165.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.165.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.166.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.166.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.166.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.167.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.167.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.167.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.168.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.168.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.168.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.169.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.169.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.169.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.170.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.170.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.170.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.171.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.171.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.171.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.172.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.172.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.172.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.173.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.173.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.173.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.174.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.174.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.174.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.175.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.175.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.175.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.176.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.176.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.176.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.177.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.177.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.177.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.178.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.178.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.178.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.179.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.179.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.179.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.180.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.180.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.180.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.181.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.181.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.181.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.182.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.182.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.182.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.183.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.183.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.183.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.184.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.184.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.184.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.185.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.185.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.185.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.186.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.186.up_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.186.down_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.187.gate_proj.weight": "model-00044-of-000163.safetensors", + "model.layers.18.mlp.experts.187.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.187.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.188.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.188.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.188.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.189.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.189.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.189.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.190.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.190.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.190.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.191.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.191.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.191.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.192.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.192.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.192.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.193.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.193.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.193.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.194.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.194.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.194.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.195.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.195.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.195.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.196.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.196.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.196.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.197.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.197.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.197.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.198.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.198.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.198.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.199.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.199.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.199.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.200.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.200.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.200.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.201.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.201.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.201.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.202.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.202.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.202.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.203.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.203.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.203.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.204.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.204.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.204.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.205.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.205.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.205.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.206.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.206.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.206.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.207.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.207.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.207.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.208.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.208.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.208.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.209.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.209.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.209.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.210.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.210.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.210.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.211.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.211.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.211.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.212.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.212.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.212.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.213.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.213.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.213.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.214.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.214.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.214.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.215.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.215.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.215.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.216.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.216.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.216.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.217.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.217.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.217.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.218.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.218.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.218.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.219.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.219.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.219.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.220.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.220.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.220.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.221.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.221.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.221.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.222.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.222.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.222.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.223.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.223.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.223.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.224.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.224.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.224.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.225.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.225.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.225.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.226.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.226.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.226.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.227.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.227.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.227.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.228.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.228.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.228.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.229.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.229.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.229.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.230.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.230.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.230.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.231.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.231.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.231.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.232.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.232.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.232.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.233.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.233.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.233.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.234.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.234.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.234.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.235.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.235.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.235.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.236.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.236.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.236.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.237.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.237.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.237.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.238.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.238.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.238.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.239.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.239.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.239.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.240.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.240.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.240.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.241.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.241.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.241.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.242.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.242.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.242.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.243.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.243.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.243.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.244.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.244.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.244.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.245.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.245.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.245.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.246.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.246.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.246.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.247.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.247.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.247.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.248.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.248.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.248.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.249.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.249.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.249.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.250.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.250.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.250.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.251.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.251.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.251.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.252.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.252.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.252.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.253.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.253.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.253.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.254.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.254.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.254.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.255.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.255.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.mlp.experts.255.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.18.input_layernorm.weight": "model-00045-of-000163.safetensors", + "model.layers.18.post_attention_layernorm.weight": "model-00045-of-000163.safetensors", + "model.layers.19.self_attn.q_a_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.self_attn.q_a_layernorm.weight": "model-00045-of-000163.safetensors", + "model.layers.19.self_attn.q_b_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.self_attn.kv_a_proj_with_mqa.weight": "model-00045-of-000163.safetensors", + "model.layers.19.self_attn.kv_a_layernorm.weight": "model-00045-of-000163.safetensors", + "model.layers.19.self_attn.kv_b_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.self_attn.o_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.gate.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.gate.e_score_correction_bias": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.shared_experts.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.shared_experts.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.shared_experts.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.0.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.0.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.0.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.1.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.1.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.1.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.2.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.2.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.2.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.3.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.3.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.3.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.4.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.4.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.4.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.5.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.5.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.5.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.6.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.6.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.6.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.7.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.7.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.7.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.8.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.8.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.8.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.9.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.9.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.9.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.10.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.10.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.10.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.11.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.11.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.11.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.12.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.12.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.12.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.13.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.13.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.13.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.14.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.14.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.14.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.15.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.15.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.15.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.16.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.16.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.16.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.17.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.17.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.17.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.18.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.18.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.18.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.19.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.19.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.19.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.20.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.20.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.20.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.21.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.21.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.21.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.22.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.22.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.22.down_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.23.gate_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.23.up_proj.weight": "model-00045-of-000163.safetensors", + "model.layers.19.mlp.experts.23.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.24.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.24.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.24.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.25.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.25.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.25.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.26.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.26.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.26.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.27.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.27.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.27.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.28.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.28.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.28.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.29.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.29.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.29.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.30.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.30.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.30.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.31.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.31.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.31.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.32.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.32.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.32.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.33.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.33.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.33.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.34.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.34.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.34.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.35.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.35.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.35.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.36.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.36.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.36.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.37.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.37.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.37.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.38.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.38.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.38.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.39.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.39.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.39.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.40.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.40.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.40.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.41.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.41.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.41.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.42.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.42.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.42.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.43.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.43.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.43.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.44.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.44.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.44.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.45.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.45.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.45.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.46.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.46.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.46.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.47.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.47.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.47.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.48.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.48.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.48.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.49.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.49.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.49.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.50.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.50.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.50.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.51.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.51.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.51.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.52.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.52.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.52.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.53.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.53.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.53.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.54.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.54.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.54.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.55.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.55.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.55.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.56.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.56.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.56.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.57.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.57.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.57.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.58.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.58.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.58.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.59.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.59.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.59.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.60.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.60.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.60.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.61.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.61.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.61.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.62.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.62.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.62.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.63.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.63.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.63.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.64.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.64.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.64.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.65.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.65.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.65.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.66.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.66.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.66.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.67.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.67.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.67.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.68.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.68.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.68.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.69.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.69.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.69.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.70.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.70.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.70.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.71.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.71.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.71.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.72.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.72.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.72.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.73.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.73.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.73.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.74.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.74.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.74.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.75.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.75.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.75.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.76.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.76.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.76.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.77.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.77.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.77.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.78.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.78.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.78.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.79.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.79.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.79.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.80.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.80.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.80.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.81.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.81.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.81.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.82.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.82.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.82.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.83.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.83.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.83.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.84.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.84.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.84.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.85.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.85.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.85.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.86.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.86.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.86.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.87.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.87.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.87.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.88.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.88.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.88.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.89.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.89.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.89.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.90.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.90.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.90.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.91.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.91.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.91.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.92.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.92.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.92.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.93.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.93.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.93.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.94.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.94.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.94.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.95.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.95.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.95.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.96.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.96.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.96.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.97.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.97.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.97.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.98.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.98.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.98.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.99.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.99.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.99.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.100.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.100.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.100.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.101.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.101.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.101.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.102.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.102.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.102.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.103.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.103.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.103.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.104.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.104.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.104.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.105.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.105.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.105.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.106.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.106.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.106.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.107.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.107.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.107.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.108.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.108.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.108.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.109.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.109.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.109.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.110.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.110.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.110.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.111.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.111.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.111.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.112.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.112.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.112.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.113.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.113.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.113.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.114.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.114.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.114.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.115.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.115.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.115.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.116.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.116.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.116.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.117.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.117.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.117.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.118.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.118.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.118.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.119.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.119.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.119.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.120.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.120.up_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.120.down_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.121.gate_proj.weight": "model-00046-of-000163.safetensors", + "model.layers.19.mlp.experts.121.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.121.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.122.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.122.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.122.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.123.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.123.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.123.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.124.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.124.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.124.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.125.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.125.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.125.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.126.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.126.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.126.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.127.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.127.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.127.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.128.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.128.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.128.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.129.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.129.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.129.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.130.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.130.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.130.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.131.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.131.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.131.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.132.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.132.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.132.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.133.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.133.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.133.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.134.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.134.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.134.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.135.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.135.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.135.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.136.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.136.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.136.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.137.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.137.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.137.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.138.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.138.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.138.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.139.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.139.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.139.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.140.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.140.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.140.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.141.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.141.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.141.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.142.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.142.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.142.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.143.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.143.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.143.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.144.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.144.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.144.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.145.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.145.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.145.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.146.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.146.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.146.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.147.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.147.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.147.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.148.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.148.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.148.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.149.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.149.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.149.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.150.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.150.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.150.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.151.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.151.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.151.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.152.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.152.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.152.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.153.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.153.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.153.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.154.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.154.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.154.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.155.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.155.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.155.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.156.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.156.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.156.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.157.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.157.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.157.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.158.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.158.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.158.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.159.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.159.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.159.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.160.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.160.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.160.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.161.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.161.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.161.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.162.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.162.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.162.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.163.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.163.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.163.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.164.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.164.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.164.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.165.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.165.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.165.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.166.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.166.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.166.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.167.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.167.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.167.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.168.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.168.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.168.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.169.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.169.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.169.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.170.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.170.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.170.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.171.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.171.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.171.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.172.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.172.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.172.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.173.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.173.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.173.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.174.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.174.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.174.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.175.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.175.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.175.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.176.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.176.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.176.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.177.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.177.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.177.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.178.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.178.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.178.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.179.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.179.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.179.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.180.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.180.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.180.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.181.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.181.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.181.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.182.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.182.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.182.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.183.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.183.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.183.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.184.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.184.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.184.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.185.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.185.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.185.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.186.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.186.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.186.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.187.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.187.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.187.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.188.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.188.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.188.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.189.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.189.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.189.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.190.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.190.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.190.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.191.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.191.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.191.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.192.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.192.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.192.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.193.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.193.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.193.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.194.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.194.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.194.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.195.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.195.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.195.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.196.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.196.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.196.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.197.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.197.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.197.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.198.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.198.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.198.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.199.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.199.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.199.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.200.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.200.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.200.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.201.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.201.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.201.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.202.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.202.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.202.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.203.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.203.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.203.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.204.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.204.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.204.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.205.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.205.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.205.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.206.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.206.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.206.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.207.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.207.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.207.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.208.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.208.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.208.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.209.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.209.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.209.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.210.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.210.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.210.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.211.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.211.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.211.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.212.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.212.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.212.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.213.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.213.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.213.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.214.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.214.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.214.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.215.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.215.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.215.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.216.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.216.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.216.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.217.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.217.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.217.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.218.gate_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.218.up_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.218.down_proj.weight": "model-00047-of-000163.safetensors", + "model.layers.19.mlp.experts.219.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.219.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.219.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.220.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.220.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.220.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.221.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.221.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.221.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.222.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.222.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.222.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.223.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.223.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.223.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.224.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.224.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.224.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.225.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.225.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.225.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.226.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.226.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.226.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.227.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.227.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.227.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.228.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.228.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.228.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.229.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.229.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.229.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.230.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.230.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.230.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.231.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.231.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.231.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.232.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.232.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.232.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.233.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.233.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.233.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.234.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.234.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.234.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.235.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.235.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.235.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.236.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.236.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.236.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.237.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.237.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.237.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.238.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.238.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.238.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.239.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.239.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.239.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.240.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.240.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.240.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.241.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.241.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.241.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.242.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.242.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.242.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.243.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.243.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.243.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.244.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.244.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.244.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.245.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.245.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.245.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.246.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.246.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.246.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.247.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.247.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.247.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.248.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.248.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.248.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.249.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.249.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.249.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.250.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.250.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.250.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.251.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.251.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.251.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.252.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.252.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.252.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.253.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.253.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.253.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.254.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.254.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.254.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.255.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.255.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.mlp.experts.255.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.19.input_layernorm.weight": "model-00048-of-000163.safetensors", + "model.layers.19.post_attention_layernorm.weight": "model-00048-of-000163.safetensors", + "model.layers.20.self_attn.q_a_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.self_attn.q_a_layernorm.weight": "model-00048-of-000163.safetensors", + "model.layers.20.self_attn.q_b_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.self_attn.kv_a_proj_with_mqa.weight": "model-00048-of-000163.safetensors", + "model.layers.20.self_attn.kv_a_layernorm.weight": "model-00048-of-000163.safetensors", + "model.layers.20.self_attn.kv_b_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.self_attn.o_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.gate.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.gate.e_score_correction_bias": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.shared_experts.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.shared_experts.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.shared_experts.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.0.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.0.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.0.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.1.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.1.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.1.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.2.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.2.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.2.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.3.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.3.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.3.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.4.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.4.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.4.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.5.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.5.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.5.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.6.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.6.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.6.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.7.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.7.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.7.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.8.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.8.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.8.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.9.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.9.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.9.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.10.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.10.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.10.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.11.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.11.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.11.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.12.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.12.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.12.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.13.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.13.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.13.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.14.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.14.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.14.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.15.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.15.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.15.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.16.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.16.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.16.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.17.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.17.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.17.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.18.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.18.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.18.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.19.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.19.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.19.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.20.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.20.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.20.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.21.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.21.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.21.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.22.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.22.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.22.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.23.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.23.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.23.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.24.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.24.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.24.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.25.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.25.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.25.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.26.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.26.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.26.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.27.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.27.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.27.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.28.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.28.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.28.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.29.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.29.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.29.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.30.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.30.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.30.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.31.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.31.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.31.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.32.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.32.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.32.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.33.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.33.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.33.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.34.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.34.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.34.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.35.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.35.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.35.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.36.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.36.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.36.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.37.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.37.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.37.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.38.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.38.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.38.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.39.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.39.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.39.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.40.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.40.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.40.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.41.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.41.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.41.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.42.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.42.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.42.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.43.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.43.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.43.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.44.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.44.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.44.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.45.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.45.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.45.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.46.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.46.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.46.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.47.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.47.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.47.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.48.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.48.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.48.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.49.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.49.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.49.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.50.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.50.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.50.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.51.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.51.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.51.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.52.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.52.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.52.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.53.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.53.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.53.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.54.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.54.up_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.54.down_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.55.gate_proj.weight": "model-00048-of-000163.safetensors", + "model.layers.20.mlp.experts.55.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.55.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.56.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.56.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.56.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.57.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.57.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.57.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.58.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.58.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.58.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.59.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.59.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.59.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.60.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.60.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.60.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.61.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.61.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.61.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.62.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.62.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.62.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.63.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.63.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.63.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.64.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.64.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.64.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.65.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.65.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.65.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.66.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.66.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.66.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.67.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.67.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.67.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.68.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.68.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.68.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.69.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.69.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.69.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.70.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.70.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.70.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.71.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.71.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.71.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.72.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.72.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.72.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.73.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.73.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.73.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.74.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.74.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.74.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.75.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.75.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.75.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.76.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.76.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.76.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.77.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.77.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.77.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.78.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.78.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.78.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.79.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.79.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.79.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.80.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.80.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.80.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.81.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.81.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.81.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.82.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.82.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.82.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.83.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.83.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.83.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.84.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.84.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.84.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.85.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.85.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.85.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.86.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.86.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.86.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.87.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.87.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.87.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.88.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.88.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.88.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.89.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.89.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.89.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.90.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.90.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.90.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.91.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.91.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.91.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.92.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.92.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.92.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.93.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.93.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.93.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.94.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.94.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.94.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.95.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.95.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.95.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.96.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.96.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.96.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.97.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.97.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.97.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.98.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.98.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.98.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.99.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.99.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.99.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.100.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.100.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.100.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.101.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.101.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.101.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.102.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.102.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.102.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.103.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.103.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.103.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.104.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.104.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.104.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.105.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.105.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.105.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.106.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.106.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.106.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.107.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.107.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.107.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.108.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.108.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.108.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.109.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.109.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.109.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.110.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.110.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.110.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.111.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.111.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.111.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.112.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.112.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.112.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.113.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.113.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.113.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.114.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.114.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.114.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.115.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.115.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.115.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.116.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.116.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.116.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.117.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.117.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.117.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.118.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.118.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.118.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.119.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.119.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.119.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.120.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.120.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.120.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.121.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.121.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.121.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.122.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.122.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.122.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.123.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.123.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.123.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.124.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.124.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.124.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.125.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.125.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.125.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.126.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.126.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.126.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.127.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.127.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.127.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.128.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.128.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.128.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.129.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.129.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.129.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.130.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.130.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.130.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.131.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.131.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.131.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.132.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.132.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.132.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.133.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.133.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.133.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.134.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.134.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.134.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.135.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.135.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.135.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.136.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.136.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.136.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.137.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.137.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.137.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.138.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.138.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.138.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.139.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.139.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.139.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.140.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.140.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.140.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.141.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.141.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.141.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.142.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.142.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.142.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.143.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.143.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.143.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.144.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.144.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.144.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.145.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.145.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.145.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.146.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.146.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.146.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.147.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.147.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.147.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.148.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.148.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.148.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.149.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.149.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.149.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.150.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.150.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.150.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.151.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.151.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.151.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.152.gate_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.152.up_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.152.down_proj.weight": "model-00049-of-000163.safetensors", + "model.layers.20.mlp.experts.153.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.153.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.153.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.154.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.154.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.154.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.155.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.155.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.155.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.156.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.156.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.156.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.157.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.157.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.157.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.158.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.158.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.158.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.159.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.159.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.159.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.160.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.160.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.160.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.161.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.161.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.161.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.162.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.162.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.162.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.163.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.163.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.163.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.164.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.164.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.164.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.165.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.165.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.165.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.166.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.166.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.166.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.167.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.167.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.167.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.168.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.168.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.168.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.169.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.169.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.169.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.170.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.170.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.170.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.171.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.171.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.171.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.172.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.172.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.172.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.173.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.173.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.173.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.174.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.174.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.174.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.175.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.175.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.175.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.176.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.176.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.176.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.177.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.177.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.177.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.178.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.178.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.178.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.179.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.179.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.179.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.180.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.180.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.180.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.181.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.181.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.181.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.182.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.182.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.182.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.183.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.183.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.183.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.184.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.184.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.184.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.185.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.185.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.185.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.186.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.186.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.186.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.187.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.187.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.187.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.188.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.188.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.188.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.189.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.189.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.189.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.190.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.190.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.190.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.191.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.191.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.191.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.192.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.192.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.192.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.193.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.193.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.193.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.194.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.194.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.194.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.195.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.195.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.195.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.196.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.196.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.196.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.197.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.197.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.197.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.198.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.198.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.198.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.199.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.199.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.199.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.200.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.200.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.200.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.201.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.201.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.201.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.202.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.202.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.202.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.203.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.203.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.203.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.204.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.204.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.204.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.205.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.205.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.205.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.206.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.206.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.206.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.207.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.207.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.207.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.208.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.208.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.208.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.209.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.209.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.209.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.210.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.210.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.210.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.211.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.211.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.211.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.212.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.212.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.212.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.213.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.213.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.213.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.214.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.214.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.214.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.215.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.215.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.215.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.216.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.216.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.216.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.217.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.217.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.217.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.218.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.218.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.218.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.219.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.219.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.219.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.220.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.220.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.220.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.221.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.221.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.221.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.222.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.222.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.222.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.223.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.223.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.223.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.224.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.224.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.224.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.225.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.225.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.225.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.226.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.226.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.226.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.227.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.227.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.227.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.228.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.228.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.228.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.229.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.229.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.229.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.230.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.230.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.230.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.231.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.231.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.231.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.232.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.232.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.232.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.233.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.233.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.233.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.234.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.234.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.234.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.235.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.235.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.235.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.236.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.236.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.236.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.237.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.237.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.237.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.238.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.238.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.238.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.239.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.239.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.239.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.240.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.240.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.240.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.241.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.241.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.241.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.242.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.242.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.242.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.243.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.243.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.243.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.244.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.244.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.244.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.245.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.245.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.245.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.246.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.246.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.246.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.247.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.247.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.247.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.248.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.248.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.248.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.249.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.249.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.249.down_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.250.gate_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.250.up_proj.weight": "model-00050-of-000163.safetensors", + "model.layers.20.mlp.experts.250.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.20.mlp.experts.251.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.20.mlp.experts.251.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.20.mlp.experts.251.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.20.mlp.experts.252.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.20.mlp.experts.252.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.20.mlp.experts.252.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.20.mlp.experts.253.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.20.mlp.experts.253.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.20.mlp.experts.253.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.20.mlp.experts.254.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.20.mlp.experts.254.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.20.mlp.experts.254.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.20.mlp.experts.255.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.20.mlp.experts.255.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.20.mlp.experts.255.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.20.input_layernorm.weight": "model-00051-of-000163.safetensors", + "model.layers.20.post_attention_layernorm.weight": "model-00051-of-000163.safetensors", + "model.layers.21.self_attn.q_a_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.self_attn.q_a_layernorm.weight": "model-00051-of-000163.safetensors", + "model.layers.21.self_attn.q_b_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.self_attn.kv_a_proj_with_mqa.weight": "model-00051-of-000163.safetensors", + "model.layers.21.self_attn.kv_a_layernorm.weight": "model-00051-of-000163.safetensors", + "model.layers.21.self_attn.kv_b_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.self_attn.o_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.gate.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.gate.e_score_correction_bias": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.shared_experts.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.shared_experts.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.shared_experts.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.0.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.0.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.0.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.1.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.1.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.1.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.2.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.2.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.2.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.3.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.3.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.3.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.4.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.4.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.4.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.5.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.5.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.5.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.6.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.6.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.6.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.7.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.7.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.7.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.8.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.8.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.8.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.9.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.9.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.9.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.10.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.10.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.10.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.11.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.11.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.11.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.12.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.12.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.12.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.13.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.13.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.13.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.14.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.14.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.14.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.15.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.15.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.15.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.16.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.16.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.16.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.17.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.17.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.17.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.18.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.18.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.18.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.19.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.19.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.19.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.20.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.20.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.20.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.21.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.21.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.21.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.22.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.22.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.22.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.23.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.23.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.23.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.24.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.24.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.24.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.25.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.25.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.25.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.26.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.26.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.26.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.27.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.27.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.27.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.28.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.28.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.28.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.29.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.29.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.29.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.30.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.30.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.30.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.31.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.31.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.31.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.32.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.32.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.32.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.33.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.33.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.33.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.34.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.34.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.34.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.35.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.35.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.35.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.36.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.36.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.36.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.37.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.37.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.37.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.38.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.38.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.38.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.39.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.39.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.39.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.40.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.40.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.40.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.41.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.41.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.41.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.42.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.42.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.42.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.43.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.43.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.43.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.44.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.44.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.44.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.45.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.45.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.45.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.46.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.46.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.46.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.47.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.47.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.47.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.48.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.48.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.48.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.49.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.49.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.49.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.50.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.50.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.50.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.51.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.51.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.51.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.52.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.52.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.52.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.53.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.53.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.53.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.54.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.54.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.54.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.55.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.55.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.55.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.56.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.56.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.56.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.57.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.57.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.57.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.58.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.58.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.58.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.59.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.59.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.59.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.60.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.60.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.60.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.61.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.61.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.61.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.62.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.62.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.62.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.63.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.63.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.63.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.64.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.64.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.64.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.65.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.65.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.65.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.66.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.66.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.66.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.67.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.67.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.67.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.68.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.68.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.68.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.69.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.69.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.69.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.70.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.70.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.70.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.71.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.71.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.71.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.72.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.72.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.72.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.73.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.73.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.73.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.74.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.74.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.74.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.75.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.75.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.75.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.76.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.76.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.76.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.77.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.77.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.77.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.78.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.78.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.78.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.79.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.79.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.79.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.80.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.80.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.80.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.81.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.81.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.81.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.82.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.82.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.82.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.83.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.83.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.83.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.84.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.84.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.84.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.85.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.85.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.85.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.86.gate_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.86.up_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.86.down_proj.weight": "model-00051-of-000163.safetensors", + "model.layers.21.mlp.experts.87.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.87.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.87.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.88.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.88.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.88.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.89.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.89.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.89.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.90.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.90.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.90.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.91.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.91.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.91.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.92.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.92.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.92.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.93.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.93.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.93.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.94.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.94.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.94.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.95.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.95.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.95.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.96.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.96.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.96.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.97.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.97.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.97.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.98.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.98.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.98.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.99.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.99.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.99.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.100.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.100.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.100.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.101.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.101.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.101.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.102.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.102.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.102.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.103.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.103.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.103.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.104.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.104.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.104.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.105.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.105.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.105.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.106.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.106.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.106.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.107.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.107.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.107.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.108.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.108.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.108.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.109.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.109.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.109.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.110.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.110.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.110.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.111.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.111.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.111.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.112.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.112.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.112.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.113.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.113.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.113.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.114.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.114.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.114.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.115.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.115.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.115.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.116.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.116.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.116.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.117.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.117.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.117.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.118.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.118.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.118.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.119.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.119.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.119.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.120.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.120.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.120.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.121.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.121.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.121.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.122.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.122.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.122.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.123.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.123.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.123.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.124.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.124.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.124.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.125.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.125.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.125.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.126.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.126.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.126.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.127.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.127.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.127.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.128.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.128.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.128.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.129.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.129.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.129.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.130.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.130.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.130.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.131.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.131.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.131.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.132.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.132.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.132.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.133.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.133.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.133.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.134.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.134.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.134.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.135.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.135.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.135.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.136.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.136.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.136.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.137.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.137.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.137.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.138.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.138.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.138.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.139.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.139.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.139.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.140.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.140.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.140.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.141.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.141.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.141.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.142.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.142.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.142.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.143.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.143.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.143.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.144.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.144.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.144.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.145.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.145.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.145.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.146.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.146.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.146.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.147.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.147.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.147.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.148.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.148.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.148.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.149.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.149.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.149.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.150.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.150.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.150.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.151.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.151.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.151.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.152.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.152.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.152.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.153.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.153.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.153.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.154.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.154.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.154.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.155.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.155.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.155.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.156.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.156.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.156.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.157.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.157.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.157.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.158.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.158.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.158.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.159.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.159.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.159.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.160.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.160.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.160.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.161.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.161.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.161.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.162.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.162.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.162.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.163.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.163.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.163.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.164.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.164.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.164.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.165.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.165.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.165.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.166.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.166.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.166.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.167.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.167.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.167.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.168.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.168.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.168.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.169.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.169.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.169.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.170.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.170.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.170.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.171.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.171.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.171.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.172.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.172.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.172.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.173.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.173.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.173.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.174.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.174.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.174.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.175.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.175.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.175.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.176.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.176.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.176.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.177.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.177.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.177.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.178.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.178.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.178.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.179.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.179.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.179.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.180.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.180.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.180.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.181.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.181.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.181.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.182.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.182.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.182.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.183.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.183.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.183.down_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.184.gate_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.184.up_proj.weight": "model-00052-of-000163.safetensors", + "model.layers.21.mlp.experts.184.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.185.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.185.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.185.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.186.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.186.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.186.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.187.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.187.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.187.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.188.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.188.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.188.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.189.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.189.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.189.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.190.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.190.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.190.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.191.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.191.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.191.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.192.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.192.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.192.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.193.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.193.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.193.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.194.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.194.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.194.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.195.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.195.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.195.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.196.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.196.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.196.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.197.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.197.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.197.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.198.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.198.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.198.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.199.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.199.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.199.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.200.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.200.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.200.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.201.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.201.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.201.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.202.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.202.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.202.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.203.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.203.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.203.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.204.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.204.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.204.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.205.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.205.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.205.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.206.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.206.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.206.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.207.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.207.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.207.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.208.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.208.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.208.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.209.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.209.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.209.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.210.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.210.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.210.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.211.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.211.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.211.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.212.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.212.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.212.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.213.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.213.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.213.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.214.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.214.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.214.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.215.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.215.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.215.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.216.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.216.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.216.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.217.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.217.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.217.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.218.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.218.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.218.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.219.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.219.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.219.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.220.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.220.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.220.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.221.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.221.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.221.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.222.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.222.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.222.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.223.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.223.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.223.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.224.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.224.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.224.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.225.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.225.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.225.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.226.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.226.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.226.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.227.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.227.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.227.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.228.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.228.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.228.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.229.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.229.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.229.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.230.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.230.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.230.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.231.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.231.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.231.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.232.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.232.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.232.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.233.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.233.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.233.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.234.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.234.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.234.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.235.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.235.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.235.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.236.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.236.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.236.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.237.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.237.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.237.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.238.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.238.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.238.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.239.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.239.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.239.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.240.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.240.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.240.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.241.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.241.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.241.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.242.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.242.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.242.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.243.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.243.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.243.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.244.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.244.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.244.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.245.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.245.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.245.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.246.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.246.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.246.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.247.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.247.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.247.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.248.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.248.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.248.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.249.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.249.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.249.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.250.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.250.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.250.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.251.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.251.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.251.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.252.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.252.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.252.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.253.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.253.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.253.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.254.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.254.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.254.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.255.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.255.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.mlp.experts.255.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.21.input_layernorm.weight": "model-00053-of-000163.safetensors", + "model.layers.21.post_attention_layernorm.weight": "model-00053-of-000163.safetensors", + "model.layers.22.self_attn.q_a_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.self_attn.q_a_layernorm.weight": "model-00053-of-000163.safetensors", + "model.layers.22.self_attn.q_b_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.self_attn.kv_a_proj_with_mqa.weight": "model-00053-of-000163.safetensors", + "model.layers.22.self_attn.kv_a_layernorm.weight": "model-00053-of-000163.safetensors", + "model.layers.22.self_attn.kv_b_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.self_attn.o_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.gate.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.gate.e_score_correction_bias": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.shared_experts.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.shared_experts.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.shared_experts.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.0.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.0.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.0.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.1.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.1.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.1.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.2.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.2.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.2.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.3.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.3.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.3.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.4.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.4.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.4.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.5.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.5.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.5.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.6.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.6.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.6.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.7.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.7.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.7.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.8.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.8.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.8.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.9.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.9.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.9.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.10.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.10.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.10.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.11.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.11.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.11.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.12.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.12.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.12.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.13.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.13.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.13.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.14.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.14.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.14.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.15.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.15.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.15.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.16.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.16.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.16.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.17.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.17.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.17.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.18.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.18.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.18.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.19.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.19.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.19.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.20.gate_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.20.up_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.20.down_proj.weight": "model-00053-of-000163.safetensors", + "model.layers.22.mlp.experts.21.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.21.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.21.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.22.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.22.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.22.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.23.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.23.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.23.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.24.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.24.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.24.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.25.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.25.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.25.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.26.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.26.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.26.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.27.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.27.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.27.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.28.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.28.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.28.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.29.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.29.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.29.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.30.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.30.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.30.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.31.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.31.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.31.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.32.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.32.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.32.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.33.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.33.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.33.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.34.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.34.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.34.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.35.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.35.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.35.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.36.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.36.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.36.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.37.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.37.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.37.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.38.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.38.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.38.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.39.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.39.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.39.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.40.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.40.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.40.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.41.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.41.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.41.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.42.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.42.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.42.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.43.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.43.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.43.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.44.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.44.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.44.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.45.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.45.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.45.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.46.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.46.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.46.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.47.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.47.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.47.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.48.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.48.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.48.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.49.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.49.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.49.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.50.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.50.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.50.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.51.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.51.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.51.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.52.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.52.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.52.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.53.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.53.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.53.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.54.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.54.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.54.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.55.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.55.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.55.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.56.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.56.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.56.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.57.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.57.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.57.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.58.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.58.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.58.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.59.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.59.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.59.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.60.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.60.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.60.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.61.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.61.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.61.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.62.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.62.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.62.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.63.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.63.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.63.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.64.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.64.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.64.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.65.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.65.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.65.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.66.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.66.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.66.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.67.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.67.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.67.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.68.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.68.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.68.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.69.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.69.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.69.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.70.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.70.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.70.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.71.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.71.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.71.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.72.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.72.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.72.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.73.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.73.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.73.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.74.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.74.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.74.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.75.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.75.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.75.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.76.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.76.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.76.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.77.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.77.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.77.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.78.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.78.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.78.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.79.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.79.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.79.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.80.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.80.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.80.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.81.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.81.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.81.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.82.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.82.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.82.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.83.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.83.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.83.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.84.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.84.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.84.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.85.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.85.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.85.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.86.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.86.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.86.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.87.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.87.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.87.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.88.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.88.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.88.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.89.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.89.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.89.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.90.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.90.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.90.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.91.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.91.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.91.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.92.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.92.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.92.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.93.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.93.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.93.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.94.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.94.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.94.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.95.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.95.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.95.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.96.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.96.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.96.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.97.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.97.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.97.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.98.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.98.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.98.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.99.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.99.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.99.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.100.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.100.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.100.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.101.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.101.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.101.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.102.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.102.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.102.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.103.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.103.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.103.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.104.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.104.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.104.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.105.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.105.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.105.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.106.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.106.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.106.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.107.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.107.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.107.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.108.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.108.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.108.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.109.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.109.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.109.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.110.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.110.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.110.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.111.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.111.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.111.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.112.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.112.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.112.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.113.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.113.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.113.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.114.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.114.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.114.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.115.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.115.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.115.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.116.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.116.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.116.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.117.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.117.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.117.down_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.118.gate_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.118.up_proj.weight": "model-00054-of-000163.safetensors", + "model.layers.22.mlp.experts.118.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.119.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.119.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.119.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.120.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.120.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.120.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.121.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.121.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.121.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.122.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.122.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.122.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.123.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.123.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.123.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.124.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.124.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.124.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.125.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.125.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.125.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.126.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.126.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.126.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.127.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.127.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.127.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.128.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.128.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.128.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.129.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.129.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.129.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.130.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.130.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.130.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.131.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.131.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.131.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.132.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.132.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.132.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.133.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.133.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.133.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.134.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.134.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.134.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.135.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.135.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.135.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.136.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.136.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.136.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.137.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.137.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.137.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.138.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.138.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.138.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.139.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.139.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.139.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.140.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.140.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.140.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.141.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.141.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.141.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.142.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.142.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.142.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.143.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.143.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.143.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.144.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.144.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.144.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.145.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.145.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.145.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.146.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.146.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.146.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.147.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.147.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.147.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.148.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.148.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.148.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.149.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.149.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.149.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.150.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.150.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.150.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.151.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.151.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.151.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.152.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.152.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.152.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.153.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.153.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.153.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.154.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.154.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.154.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.155.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.155.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.155.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.156.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.156.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.156.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.157.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.157.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.157.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.158.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.158.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.158.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.159.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.159.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.159.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.160.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.160.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.160.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.161.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.161.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.161.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.162.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.162.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.162.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.163.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.163.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.163.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.164.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.164.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.164.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.165.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.165.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.165.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.166.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.166.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.166.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.167.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.167.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.167.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.168.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.168.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.168.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.169.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.169.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.169.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.170.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.170.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.170.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.171.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.171.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.171.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.172.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.172.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.172.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.173.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.173.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.173.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.174.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.174.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.174.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.175.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.175.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.175.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.176.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.176.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.176.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.177.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.177.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.177.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.178.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.178.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.178.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.179.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.179.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.179.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.180.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.180.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.180.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.181.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.181.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.181.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.182.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.182.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.182.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.183.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.183.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.183.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.184.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.184.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.184.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.185.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.185.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.185.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.186.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.186.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.186.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.187.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.187.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.187.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.188.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.188.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.188.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.189.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.189.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.189.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.190.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.190.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.190.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.191.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.191.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.191.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.192.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.192.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.192.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.193.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.193.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.193.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.194.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.194.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.194.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.195.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.195.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.195.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.196.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.196.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.196.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.197.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.197.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.197.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.198.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.198.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.198.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.199.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.199.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.199.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.200.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.200.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.200.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.201.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.201.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.201.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.202.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.202.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.202.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.203.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.203.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.203.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.204.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.204.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.204.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.205.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.205.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.205.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.206.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.206.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.206.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.207.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.207.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.207.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.208.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.208.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.208.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.209.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.209.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.209.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.210.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.210.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.210.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.211.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.211.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.211.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.212.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.212.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.212.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.213.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.213.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.213.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.214.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.214.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.214.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.215.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.215.up_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.215.down_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.216.gate_proj.weight": "model-00055-of-000163.safetensors", + "model.layers.22.mlp.experts.216.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.216.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.217.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.217.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.217.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.218.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.218.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.218.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.219.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.219.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.219.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.220.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.220.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.220.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.221.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.221.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.221.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.222.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.222.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.222.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.223.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.223.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.223.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.224.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.224.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.224.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.225.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.225.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.225.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.226.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.226.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.226.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.227.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.227.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.227.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.228.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.228.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.228.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.229.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.229.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.229.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.230.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.230.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.230.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.231.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.231.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.231.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.232.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.232.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.232.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.233.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.233.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.233.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.234.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.234.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.234.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.235.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.235.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.235.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.236.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.236.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.236.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.237.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.237.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.237.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.238.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.238.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.238.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.239.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.239.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.239.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.240.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.240.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.240.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.241.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.241.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.241.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.242.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.242.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.242.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.243.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.243.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.243.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.244.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.244.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.244.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.245.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.245.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.245.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.246.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.246.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.246.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.247.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.247.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.247.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.248.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.248.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.248.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.249.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.249.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.249.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.250.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.250.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.250.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.251.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.251.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.251.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.252.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.252.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.252.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.253.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.253.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.253.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.254.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.254.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.254.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.255.gate_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.255.up_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.mlp.experts.255.down_proj.weight": "model-00056-of-000163.safetensors", + "model.layers.22.input_layernorm.weight": "model-00056-of-000163.safetensors", + "model.layers.22.post_attention_layernorm.weight": "model-00056-of-000163.safetensors", + "model.layers.23.self_attn.q_a_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.self_attn.q_a_layernorm.weight": "model-00057-of-000163.safetensors", + "model.layers.23.self_attn.q_b_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.self_attn.kv_a_proj_with_mqa.weight": "model-00057-of-000163.safetensors", + "model.layers.23.self_attn.kv_a_layernorm.weight": "model-00057-of-000163.safetensors", + "model.layers.23.self_attn.kv_b_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.self_attn.o_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.gate.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.gate.e_score_correction_bias": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.shared_experts.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.shared_experts.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.shared_experts.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.0.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.0.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.0.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.1.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.1.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.1.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.2.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.2.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.2.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.3.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.3.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.3.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.4.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.4.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.4.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.5.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.5.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.5.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.6.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.6.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.6.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.7.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.7.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.7.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.8.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.8.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.8.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.9.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.9.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.9.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.10.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.10.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.10.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.11.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.11.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.11.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.12.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.12.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.12.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.13.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.13.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.13.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.14.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.14.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.14.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.15.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.15.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.15.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.16.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.16.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.16.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.17.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.17.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.17.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.18.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.18.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.18.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.19.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.19.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.19.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.20.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.20.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.20.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.21.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.21.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.21.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.22.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.22.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.22.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.23.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.23.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.23.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.24.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.24.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.24.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.25.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.25.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.25.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.26.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.26.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.26.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.27.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.27.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.27.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.28.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.28.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.28.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.29.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.29.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.29.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.30.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.30.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.30.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.31.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.31.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.31.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.32.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.32.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.32.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.33.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.33.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.33.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.34.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.34.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.34.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.35.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.35.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.35.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.36.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.36.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.36.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.37.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.37.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.37.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.38.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.38.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.38.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.39.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.39.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.39.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.40.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.40.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.40.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.41.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.41.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.41.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.42.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.42.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.42.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.43.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.43.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.43.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.44.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.44.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.44.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.45.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.45.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.45.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.46.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.46.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.46.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.47.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.47.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.47.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.48.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.48.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.48.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.49.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.49.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.49.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.50.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.50.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.50.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.51.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.51.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.51.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.52.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.52.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.52.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.53.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.53.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.53.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.54.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.54.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.54.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.55.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.55.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.55.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.56.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.56.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.56.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.57.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.57.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.57.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.58.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.58.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.58.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.59.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.59.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.59.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.60.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.60.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.60.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.61.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.61.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.61.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.62.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.62.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.62.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.63.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.63.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.63.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.64.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.64.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.64.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.65.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.65.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.65.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.66.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.66.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.66.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.67.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.67.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.67.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.68.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.68.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.68.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.69.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.69.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.69.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.70.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.70.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.70.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.71.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.71.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.71.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.72.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.72.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.72.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.73.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.73.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.73.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.74.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.74.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.74.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.75.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.75.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.75.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.76.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.76.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.76.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.77.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.77.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.77.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.78.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.78.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.78.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.79.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.79.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.79.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.80.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.80.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.80.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.81.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.81.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.81.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.82.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.82.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.82.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.83.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.83.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.83.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.84.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.84.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.84.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.85.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.85.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.85.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.86.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.86.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.86.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.87.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.87.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.87.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.88.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.88.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.88.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.89.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.89.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.89.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.90.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.90.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.90.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.91.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.91.up_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.91.down_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.92.gate_proj.weight": "model-00057-of-000163.safetensors", + "model.layers.23.mlp.experts.92.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.92.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.93.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.93.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.93.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.94.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.94.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.94.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.95.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.95.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.95.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.96.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.96.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.96.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.97.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.97.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.97.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.98.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.98.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.98.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.99.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.99.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.99.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.100.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.100.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.100.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.101.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.101.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.101.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.102.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.102.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.102.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.103.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.103.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.103.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.104.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.104.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.104.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.105.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.105.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.105.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.106.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.106.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.106.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.107.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.107.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.107.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.108.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.108.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.108.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.109.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.109.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.109.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.110.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.110.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.110.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.111.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.111.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.111.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.112.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.112.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.112.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.113.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.113.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.113.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.114.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.114.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.114.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.115.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.115.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.115.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.116.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.116.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.116.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.117.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.117.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.117.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.118.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.118.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.118.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.119.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.119.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.119.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.120.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.120.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.120.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.121.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.121.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.121.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.122.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.122.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.122.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.123.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.123.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.123.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.124.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.124.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.124.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.125.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.125.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.125.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.126.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.126.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.126.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.127.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.127.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.127.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.128.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.128.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.128.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.129.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.129.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.129.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.130.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.130.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.130.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.131.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.131.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.131.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.132.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.132.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.132.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.133.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.133.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.133.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.134.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.134.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.134.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.135.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.135.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.135.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.136.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.136.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.136.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.137.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.137.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.137.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.138.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.138.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.138.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.139.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.139.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.139.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.140.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.140.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.140.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.141.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.141.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.141.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.142.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.142.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.142.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.143.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.143.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.143.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.144.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.144.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.144.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.145.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.145.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.145.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.146.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.146.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.146.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.147.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.147.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.147.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.148.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.148.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.148.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.149.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.149.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.149.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.150.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.150.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.150.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.151.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.151.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.151.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.152.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.152.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.152.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.153.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.153.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.153.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.154.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.154.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.154.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.155.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.155.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.155.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.156.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.156.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.156.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.157.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.157.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.157.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.158.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.158.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.158.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.159.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.159.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.159.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.160.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.160.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.160.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.161.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.161.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.161.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.162.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.162.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.162.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.163.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.163.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.163.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.164.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.164.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.164.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.165.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.165.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.165.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.166.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.166.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.166.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.167.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.167.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.167.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.168.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.168.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.168.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.169.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.169.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.169.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.170.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.170.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.170.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.171.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.171.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.171.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.172.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.172.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.172.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.173.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.173.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.173.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.174.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.174.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.174.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.175.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.175.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.175.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.176.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.176.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.176.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.177.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.177.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.177.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.178.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.178.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.178.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.179.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.179.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.179.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.180.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.180.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.180.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.181.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.181.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.181.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.182.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.182.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.182.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.183.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.183.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.183.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.184.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.184.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.184.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.185.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.185.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.185.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.186.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.186.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.186.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.187.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.187.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.187.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.188.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.188.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.188.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.189.gate_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.189.up_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.189.down_proj.weight": "model-00058-of-000163.safetensors", + "model.layers.23.mlp.experts.190.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.190.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.190.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.191.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.191.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.191.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.192.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.192.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.192.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.193.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.193.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.193.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.194.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.194.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.194.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.195.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.195.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.195.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.196.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.196.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.196.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.197.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.197.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.197.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.198.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.198.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.198.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.199.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.199.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.199.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.200.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.200.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.200.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.201.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.201.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.201.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.202.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.202.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.202.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.203.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.203.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.203.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.204.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.204.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.204.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.205.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.205.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.205.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.206.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.206.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.206.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.207.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.207.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.207.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.208.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.208.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.208.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.209.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.209.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.209.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.210.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.210.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.210.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.211.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.211.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.211.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.212.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.212.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.212.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.213.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.213.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.213.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.214.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.214.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.214.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.215.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.215.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.215.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.216.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.216.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.216.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.217.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.217.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.217.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.218.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.218.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.218.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.219.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.219.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.219.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.220.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.220.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.220.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.221.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.221.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.221.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.222.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.222.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.222.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.223.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.223.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.223.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.224.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.224.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.224.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.225.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.225.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.225.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.226.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.226.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.226.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.227.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.227.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.227.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.228.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.228.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.228.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.229.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.229.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.229.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.230.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.230.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.230.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.231.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.231.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.231.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.232.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.232.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.232.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.233.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.233.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.233.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.234.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.234.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.234.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.235.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.235.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.235.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.236.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.236.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.236.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.237.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.237.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.237.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.238.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.238.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.238.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.239.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.239.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.239.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.240.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.240.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.240.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.241.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.241.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.241.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.242.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.242.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.242.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.243.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.243.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.243.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.244.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.244.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.244.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.245.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.245.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.245.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.246.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.246.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.246.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.247.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.247.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.247.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.248.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.248.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.248.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.249.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.249.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.249.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.250.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.250.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.250.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.251.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.251.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.251.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.252.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.252.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.252.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.253.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.253.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.253.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.254.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.254.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.254.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.255.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.255.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.mlp.experts.255.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.23.input_layernorm.weight": "model-00059-of-000163.safetensors", + "model.layers.23.post_attention_layernorm.weight": "model-00059-of-000163.safetensors", + "model.layers.24.self_attn.q_a_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.self_attn.q_a_layernorm.weight": "model-00059-of-000163.safetensors", + "model.layers.24.self_attn.q_b_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.self_attn.kv_a_proj_with_mqa.weight": "model-00059-of-000163.safetensors", + "model.layers.24.self_attn.kv_a_layernorm.weight": "model-00059-of-000163.safetensors", + "model.layers.24.self_attn.kv_b_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.self_attn.o_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.gate.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.gate.e_score_correction_bias": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.shared_experts.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.shared_experts.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.shared_experts.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.0.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.0.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.0.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.1.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.1.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.1.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.2.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.2.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.2.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.3.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.3.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.3.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.4.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.4.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.4.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.5.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.5.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.5.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.6.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.6.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.6.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.7.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.7.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.7.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.8.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.8.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.8.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.9.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.9.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.9.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.10.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.10.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.10.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.11.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.11.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.11.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.12.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.12.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.12.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.13.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.13.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.13.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.14.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.14.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.14.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.15.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.15.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.15.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.16.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.16.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.16.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.17.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.17.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.17.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.18.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.18.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.18.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.19.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.19.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.19.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.20.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.20.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.20.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.21.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.21.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.21.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.22.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.22.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.22.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.23.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.23.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.23.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.24.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.24.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.24.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.25.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.25.up_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.25.down_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.26.gate_proj.weight": "model-00059-of-000163.safetensors", + "model.layers.24.mlp.experts.26.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.26.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.27.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.27.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.27.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.28.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.28.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.28.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.29.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.29.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.29.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.30.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.30.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.30.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.31.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.31.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.31.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.32.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.32.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.32.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.33.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.33.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.33.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.34.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.34.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.34.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.35.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.35.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.35.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.36.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.36.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.36.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.37.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.37.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.37.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.38.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.38.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.38.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.39.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.39.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.39.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.40.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.40.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.40.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.41.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.41.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.41.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.42.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.42.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.42.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.43.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.43.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.43.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.44.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.44.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.44.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.45.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.45.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.45.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.46.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.46.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.46.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.47.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.47.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.47.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.48.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.48.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.48.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.49.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.49.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.49.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.50.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.50.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.50.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.51.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.51.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.51.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.52.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.52.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.52.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.53.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.53.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.53.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.54.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.54.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.54.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.55.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.55.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.55.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.56.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.56.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.56.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.57.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.57.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.57.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.58.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.58.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.58.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.59.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.59.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.59.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.60.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.60.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.60.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.61.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.61.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.61.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.62.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.62.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.62.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.63.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.63.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.63.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.64.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.64.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.64.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.65.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.65.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.65.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.66.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.66.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.66.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.67.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.67.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.67.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.68.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.68.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.68.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.69.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.69.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.69.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.70.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.70.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.70.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.71.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.71.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.71.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.72.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.72.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.72.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.73.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.73.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.73.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.74.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.74.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.74.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.75.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.75.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.75.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.76.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.76.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.76.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.77.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.77.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.77.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.78.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.78.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.78.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.79.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.79.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.79.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.80.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.80.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.80.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.81.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.81.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.81.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.82.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.82.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.82.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.83.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.83.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.83.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.84.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.84.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.84.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.85.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.85.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.85.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.86.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.86.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.86.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.87.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.87.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.87.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.88.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.88.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.88.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.89.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.89.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.89.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.90.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.90.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.90.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.91.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.91.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.91.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.92.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.92.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.92.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.93.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.93.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.93.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.94.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.94.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.94.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.95.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.95.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.95.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.96.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.96.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.96.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.97.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.97.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.97.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.98.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.98.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.98.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.99.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.99.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.99.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.100.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.100.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.100.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.101.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.101.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.101.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.102.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.102.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.102.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.103.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.103.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.103.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.104.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.104.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.104.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.105.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.105.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.105.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.106.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.106.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.106.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.107.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.107.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.107.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.108.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.108.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.108.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.109.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.109.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.109.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.110.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.110.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.110.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.111.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.111.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.111.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.112.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.112.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.112.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.113.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.113.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.113.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.114.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.114.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.114.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.115.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.115.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.115.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.116.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.116.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.116.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.117.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.117.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.117.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.118.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.118.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.118.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.119.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.119.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.119.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.120.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.120.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.120.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.121.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.121.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.121.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.122.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.122.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.122.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.123.gate_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.123.up_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.123.down_proj.weight": "model-00060-of-000163.safetensors", + "model.layers.24.mlp.experts.124.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.124.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.124.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.125.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.125.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.125.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.126.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.126.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.126.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.127.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.127.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.127.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.128.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.128.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.128.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.129.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.129.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.129.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.130.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.130.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.130.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.131.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.131.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.131.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.132.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.132.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.132.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.133.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.133.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.133.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.134.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.134.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.134.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.135.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.135.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.135.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.136.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.136.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.136.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.137.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.137.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.137.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.138.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.138.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.138.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.139.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.139.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.139.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.140.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.140.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.140.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.141.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.141.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.141.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.142.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.142.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.142.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.143.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.143.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.143.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.144.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.144.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.144.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.145.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.145.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.145.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.146.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.146.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.146.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.147.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.147.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.147.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.148.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.148.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.148.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.149.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.149.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.149.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.150.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.150.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.150.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.151.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.151.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.151.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.152.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.152.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.152.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.153.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.153.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.153.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.154.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.154.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.154.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.155.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.155.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.155.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.156.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.156.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.156.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.157.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.157.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.157.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.158.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.158.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.158.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.159.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.159.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.159.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.160.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.160.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.160.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.161.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.161.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.161.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.162.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.162.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.162.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.163.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.163.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.163.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.164.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.164.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.164.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.165.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.165.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.165.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.166.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.166.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.166.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.167.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.167.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.167.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.168.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.168.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.168.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.169.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.169.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.169.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.170.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.170.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.170.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.171.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.171.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.171.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.172.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.172.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.172.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.173.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.173.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.173.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.174.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.174.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.174.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.175.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.175.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.175.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.176.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.176.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.176.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.177.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.177.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.177.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.178.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.178.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.178.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.179.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.179.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.179.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.180.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.180.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.180.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.181.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.181.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.181.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.182.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.182.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.182.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.183.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.183.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.183.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.184.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.184.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.184.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.185.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.185.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.185.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.186.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.186.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.186.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.187.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.187.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.187.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.188.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.188.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.188.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.189.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.189.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.189.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.190.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.190.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.190.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.191.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.191.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.191.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.192.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.192.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.192.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.193.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.193.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.193.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.194.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.194.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.194.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.195.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.195.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.195.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.196.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.196.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.196.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.197.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.197.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.197.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.198.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.198.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.198.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.199.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.199.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.199.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.200.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.200.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.200.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.201.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.201.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.201.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.202.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.202.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.202.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.203.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.203.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.203.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.204.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.204.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.204.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.205.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.205.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.205.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.206.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.206.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.206.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.207.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.207.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.207.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.208.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.208.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.208.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.209.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.209.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.209.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.210.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.210.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.210.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.211.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.211.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.211.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.212.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.212.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.212.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.213.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.213.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.213.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.214.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.214.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.214.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.215.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.215.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.215.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.216.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.216.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.216.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.217.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.217.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.217.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.218.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.218.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.218.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.219.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.219.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.219.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.220.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.220.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.220.down_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.221.gate_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.221.up_proj.weight": "model-00061-of-000163.safetensors", + "model.layers.24.mlp.experts.221.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.222.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.222.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.222.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.223.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.223.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.223.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.224.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.224.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.224.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.225.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.225.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.225.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.226.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.226.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.226.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.227.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.227.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.227.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.228.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.228.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.228.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.229.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.229.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.229.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.230.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.230.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.230.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.231.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.231.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.231.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.232.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.232.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.232.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.233.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.233.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.233.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.234.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.234.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.234.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.235.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.235.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.235.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.236.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.236.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.236.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.237.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.237.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.237.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.238.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.238.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.238.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.239.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.239.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.239.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.240.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.240.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.240.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.241.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.241.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.241.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.242.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.242.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.242.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.243.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.243.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.243.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.244.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.244.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.244.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.245.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.245.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.245.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.246.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.246.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.246.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.247.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.247.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.247.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.248.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.248.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.248.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.249.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.249.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.249.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.250.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.250.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.250.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.251.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.251.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.251.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.252.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.252.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.252.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.253.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.253.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.253.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.254.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.254.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.254.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.255.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.255.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.mlp.experts.255.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.24.input_layernorm.weight": "model-00062-of-000163.safetensors", + "model.layers.24.post_attention_layernorm.weight": "model-00062-of-000163.safetensors", + "model.layers.25.self_attn.q_a_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.self_attn.q_a_layernorm.weight": "model-00062-of-000163.safetensors", + "model.layers.25.self_attn.q_b_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.self_attn.kv_a_proj_with_mqa.weight": "model-00062-of-000163.safetensors", + "model.layers.25.self_attn.kv_a_layernorm.weight": "model-00062-of-000163.safetensors", + "model.layers.25.self_attn.kv_b_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.self_attn.o_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.gate.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.gate.e_score_correction_bias": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.shared_experts.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.shared_experts.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.shared_experts.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.0.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.0.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.0.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.1.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.1.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.1.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.2.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.2.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.2.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.3.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.3.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.3.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.4.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.4.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.4.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.5.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.5.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.5.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.6.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.6.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.6.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.7.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.7.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.7.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.8.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.8.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.8.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.9.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.9.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.9.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.10.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.10.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.10.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.11.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.11.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.11.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.12.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.12.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.12.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.13.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.13.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.13.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.14.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.14.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.14.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.15.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.15.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.15.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.16.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.16.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.16.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.17.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.17.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.17.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.18.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.18.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.18.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.19.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.19.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.19.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.20.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.20.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.20.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.21.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.21.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.21.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.22.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.22.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.22.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.23.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.23.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.23.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.24.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.24.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.24.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.25.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.25.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.25.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.26.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.26.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.26.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.27.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.27.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.27.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.28.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.28.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.28.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.29.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.29.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.29.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.30.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.30.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.30.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.31.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.31.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.31.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.32.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.32.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.32.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.33.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.33.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.33.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.34.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.34.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.34.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.35.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.35.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.35.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.36.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.36.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.36.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.37.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.37.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.37.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.38.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.38.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.38.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.39.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.39.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.39.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.40.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.40.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.40.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.41.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.41.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.41.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.42.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.42.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.42.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.43.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.43.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.43.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.44.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.44.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.44.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.45.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.45.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.45.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.46.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.46.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.46.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.47.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.47.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.47.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.48.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.48.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.48.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.49.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.49.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.49.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.50.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.50.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.50.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.51.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.51.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.51.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.52.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.52.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.52.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.53.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.53.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.53.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.54.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.54.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.54.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.55.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.55.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.55.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.56.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.56.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.56.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.57.gate_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.57.up_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.57.down_proj.weight": "model-00062-of-000163.safetensors", + "model.layers.25.mlp.experts.58.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.58.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.58.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.59.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.59.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.59.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.60.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.60.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.60.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.61.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.61.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.61.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.62.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.62.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.62.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.63.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.63.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.63.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.64.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.64.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.64.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.65.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.65.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.65.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.66.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.66.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.66.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.67.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.67.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.67.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.68.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.68.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.68.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.69.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.69.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.69.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.70.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.70.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.70.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.71.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.71.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.71.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.72.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.72.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.72.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.73.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.73.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.73.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.74.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.74.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.74.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.75.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.75.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.75.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.76.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.76.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.76.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.77.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.77.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.77.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.78.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.78.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.78.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.79.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.79.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.79.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.80.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.80.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.80.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.81.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.81.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.81.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.82.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.82.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.82.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.83.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.83.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.83.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.84.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.84.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.84.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.85.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.85.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.85.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.86.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.86.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.86.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.87.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.87.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.87.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.88.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.88.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.88.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.89.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.89.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.89.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.90.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.90.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.90.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.91.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.91.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.91.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.92.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.92.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.92.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.93.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.93.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.93.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.94.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.94.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.94.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.95.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.95.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.95.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.96.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.96.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.96.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.97.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.97.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.97.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.98.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.98.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.98.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.99.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.99.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.99.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.100.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.100.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.100.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.101.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.101.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.101.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.102.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.102.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.102.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.103.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.103.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.103.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.104.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.104.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.104.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.105.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.105.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.105.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.106.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.106.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.106.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.107.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.107.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.107.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.108.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.108.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.108.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.109.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.109.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.109.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.110.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.110.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.110.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.111.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.111.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.111.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.112.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.112.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.112.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.113.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.113.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.113.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.114.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.114.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.114.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.115.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.115.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.115.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.116.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.116.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.116.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.117.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.117.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.117.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.118.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.118.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.118.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.119.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.119.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.119.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.120.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.120.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.120.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.121.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.121.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.121.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.122.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.122.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.122.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.123.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.123.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.123.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.124.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.124.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.124.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.125.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.125.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.125.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.126.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.126.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.126.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.127.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.127.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.127.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.128.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.128.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.128.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.129.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.129.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.129.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.130.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.130.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.130.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.131.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.131.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.131.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.132.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.132.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.132.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.133.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.133.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.133.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.134.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.134.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.134.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.135.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.135.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.135.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.136.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.136.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.136.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.137.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.137.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.137.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.138.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.138.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.138.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.139.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.139.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.139.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.140.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.140.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.140.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.141.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.141.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.141.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.142.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.142.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.142.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.143.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.143.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.143.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.144.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.144.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.144.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.145.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.145.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.145.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.146.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.146.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.146.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.147.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.147.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.147.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.148.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.148.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.148.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.149.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.149.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.149.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.150.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.150.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.150.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.151.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.151.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.151.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.152.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.152.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.152.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.153.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.153.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.153.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.154.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.154.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.154.down_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.155.gate_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.155.up_proj.weight": "model-00063-of-000163.safetensors", + "model.layers.25.mlp.experts.155.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.156.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.156.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.156.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.157.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.157.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.157.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.158.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.158.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.158.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.159.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.159.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.159.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.160.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.160.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.160.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.161.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.161.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.161.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.162.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.162.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.162.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.163.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.163.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.163.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.164.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.164.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.164.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.165.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.165.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.165.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.166.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.166.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.166.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.167.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.167.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.167.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.168.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.168.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.168.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.169.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.169.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.169.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.170.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.170.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.170.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.171.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.171.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.171.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.172.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.172.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.172.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.173.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.173.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.173.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.174.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.174.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.174.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.175.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.175.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.175.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.176.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.176.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.176.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.177.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.177.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.177.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.178.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.178.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.178.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.179.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.179.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.179.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.180.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.180.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.180.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.181.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.181.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.181.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.182.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.182.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.182.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.183.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.183.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.183.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.184.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.184.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.184.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.185.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.185.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.185.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.186.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.186.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.186.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.187.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.187.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.187.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.188.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.188.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.188.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.189.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.189.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.189.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.190.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.190.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.190.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.191.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.191.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.191.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.192.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.192.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.192.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.193.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.193.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.193.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.194.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.194.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.194.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.195.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.195.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.195.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.196.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.196.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.196.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.197.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.197.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.197.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.198.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.198.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.198.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.199.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.199.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.199.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.200.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.200.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.200.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.201.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.201.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.201.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.202.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.202.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.202.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.203.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.203.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.203.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.204.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.204.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.204.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.205.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.205.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.205.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.206.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.206.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.206.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.207.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.207.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.207.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.208.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.208.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.208.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.209.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.209.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.209.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.210.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.210.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.210.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.211.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.211.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.211.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.212.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.212.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.212.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.213.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.213.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.213.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.214.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.214.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.214.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.215.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.215.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.215.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.216.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.216.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.216.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.217.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.217.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.217.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.218.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.218.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.218.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.219.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.219.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.219.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.220.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.220.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.220.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.221.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.221.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.221.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.222.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.222.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.222.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.223.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.223.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.223.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.224.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.224.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.224.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.225.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.225.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.225.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.226.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.226.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.226.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.227.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.227.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.227.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.228.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.228.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.228.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.229.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.229.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.229.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.230.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.230.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.230.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.231.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.231.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.231.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.232.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.232.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.232.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.233.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.233.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.233.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.234.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.234.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.234.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.235.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.235.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.235.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.236.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.236.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.236.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.237.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.237.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.237.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.238.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.238.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.238.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.239.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.239.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.239.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.240.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.240.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.240.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.241.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.241.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.241.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.242.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.242.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.242.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.243.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.243.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.243.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.244.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.244.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.244.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.245.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.245.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.245.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.246.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.246.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.246.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.247.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.247.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.247.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.248.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.248.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.248.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.249.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.249.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.249.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.250.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.250.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.250.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.251.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.251.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.251.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.252.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.252.up_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.252.down_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.253.gate_proj.weight": "model-00064-of-000163.safetensors", + "model.layers.25.mlp.experts.253.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.25.mlp.experts.253.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.25.mlp.experts.254.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.25.mlp.experts.254.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.25.mlp.experts.254.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.25.mlp.experts.255.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.25.mlp.experts.255.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.25.mlp.experts.255.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.25.input_layernorm.weight": "model-00065-of-000163.safetensors", + "model.layers.25.post_attention_layernorm.weight": "model-00065-of-000163.safetensors", + "model.layers.26.self_attn.q_a_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.self_attn.q_a_layernorm.weight": "model-00065-of-000163.safetensors", + "model.layers.26.self_attn.q_b_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.self_attn.kv_a_proj_with_mqa.weight": "model-00065-of-000163.safetensors", + "model.layers.26.self_attn.kv_a_layernorm.weight": "model-00065-of-000163.safetensors", + "model.layers.26.self_attn.kv_b_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.self_attn.o_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.gate.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.gate.e_score_correction_bias": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.shared_experts.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.shared_experts.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.shared_experts.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.0.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.0.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.0.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.1.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.1.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.1.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.2.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.2.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.2.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.3.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.3.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.3.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.4.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.4.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.4.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.5.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.5.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.5.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.6.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.6.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.6.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.7.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.7.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.7.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.8.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.8.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.8.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.9.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.9.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.9.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.10.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.10.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.10.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.11.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.11.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.11.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.12.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.12.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.12.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.13.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.13.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.13.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.14.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.14.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.14.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.15.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.15.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.15.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.16.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.16.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.16.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.17.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.17.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.17.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.18.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.18.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.18.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.19.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.19.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.19.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.20.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.20.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.20.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.21.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.21.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.21.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.22.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.22.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.22.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.23.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.23.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.23.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.24.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.24.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.24.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.25.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.25.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.25.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.26.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.26.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.26.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.27.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.27.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.27.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.28.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.28.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.28.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.29.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.29.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.29.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.30.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.30.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.30.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.31.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.31.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.31.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.32.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.32.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.32.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.33.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.33.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.33.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.34.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.34.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.34.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.35.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.35.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.35.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.36.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.36.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.36.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.37.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.37.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.37.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.38.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.38.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.38.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.39.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.39.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.39.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.40.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.40.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.40.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.41.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.41.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.41.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.42.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.42.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.42.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.43.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.43.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.43.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.44.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.44.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.44.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.45.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.45.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.45.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.46.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.46.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.46.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.47.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.47.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.47.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.48.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.48.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.48.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.49.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.49.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.49.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.50.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.50.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.50.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.51.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.51.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.51.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.52.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.52.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.52.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.53.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.53.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.53.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.54.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.54.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.54.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.55.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.55.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.55.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.56.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.56.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.56.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.57.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.57.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.57.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.58.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.58.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.58.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.59.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.59.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.59.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.60.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.60.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.60.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.61.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.61.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.61.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.62.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.62.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.62.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.63.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.63.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.63.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.64.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.64.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.64.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.65.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.65.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.65.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.66.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.66.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.66.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.67.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.67.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.67.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.68.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.68.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.68.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.69.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.69.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.69.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.70.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.70.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.70.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.71.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.71.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.71.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.72.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.72.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.72.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.73.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.73.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.73.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.74.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.74.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.74.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.75.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.75.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.75.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.76.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.76.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.76.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.77.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.77.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.77.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.78.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.78.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.78.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.79.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.79.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.79.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.80.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.80.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.80.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.81.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.81.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.81.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.82.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.82.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.82.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.83.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.83.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.83.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.84.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.84.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.84.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.85.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.85.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.85.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.86.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.86.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.86.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.87.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.87.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.87.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.88.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.88.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.88.down_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.89.gate_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.89.up_proj.weight": "model-00065-of-000163.safetensors", + "model.layers.26.mlp.experts.89.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.90.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.90.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.90.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.91.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.91.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.91.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.92.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.92.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.92.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.93.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.93.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.93.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.94.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.94.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.94.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.95.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.95.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.95.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.96.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.96.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.96.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.97.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.97.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.97.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.98.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.98.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.98.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.99.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.99.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.99.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.100.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.100.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.100.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.101.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.101.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.101.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.102.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.102.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.102.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.103.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.103.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.103.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.104.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.104.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.104.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.105.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.105.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.105.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.106.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.106.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.106.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.107.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.107.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.107.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.108.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.108.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.108.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.109.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.109.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.109.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.110.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.110.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.110.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.111.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.111.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.111.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.112.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.112.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.112.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.113.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.113.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.113.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.114.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.114.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.114.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.115.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.115.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.115.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.116.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.116.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.116.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.117.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.117.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.117.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.118.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.118.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.118.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.119.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.119.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.119.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.120.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.120.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.120.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.121.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.121.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.121.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.122.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.122.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.122.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.123.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.123.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.123.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.124.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.124.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.124.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.125.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.125.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.125.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.126.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.126.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.126.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.127.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.127.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.127.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.128.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.128.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.128.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.129.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.129.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.129.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.130.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.130.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.130.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.131.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.131.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.131.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.132.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.132.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.132.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.133.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.133.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.133.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.134.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.134.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.134.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.135.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.135.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.135.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.136.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.136.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.136.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.137.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.137.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.137.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.138.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.138.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.138.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.139.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.139.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.139.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.140.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.140.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.140.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.141.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.141.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.141.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.142.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.142.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.142.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.143.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.143.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.143.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.144.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.144.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.144.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.145.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.145.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.145.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.146.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.146.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.146.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.147.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.147.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.147.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.148.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.148.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.148.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.149.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.149.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.149.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.150.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.150.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.150.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.151.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.151.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.151.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.152.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.152.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.152.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.153.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.153.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.153.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.154.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.154.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.154.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.155.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.155.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.155.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.156.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.156.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.156.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.157.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.157.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.157.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.158.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.158.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.158.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.159.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.159.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.159.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.160.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.160.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.160.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.161.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.161.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.161.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.162.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.162.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.162.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.163.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.163.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.163.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.164.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.164.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.164.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.165.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.165.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.165.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.166.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.166.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.166.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.167.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.167.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.167.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.168.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.168.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.168.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.169.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.169.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.169.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.170.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.170.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.170.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.171.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.171.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.171.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.172.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.172.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.172.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.173.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.173.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.173.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.174.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.174.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.174.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.175.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.175.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.175.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.176.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.176.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.176.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.177.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.177.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.177.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.178.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.178.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.178.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.179.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.179.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.179.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.180.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.180.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.180.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.181.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.181.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.181.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.182.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.182.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.182.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.183.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.183.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.183.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.184.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.184.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.184.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.185.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.185.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.185.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.186.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.186.up_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.186.down_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.187.gate_proj.weight": "model-00066-of-000163.safetensors", + "model.layers.26.mlp.experts.187.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.187.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.188.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.188.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.188.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.189.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.189.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.189.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.190.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.190.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.190.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.191.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.191.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.191.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.192.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.192.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.192.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.193.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.193.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.193.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.194.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.194.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.194.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.195.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.195.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.195.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.196.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.196.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.196.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.197.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.197.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.197.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.198.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.198.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.198.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.199.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.199.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.199.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.200.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.200.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.200.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.201.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.201.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.201.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.202.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.202.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.202.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.203.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.203.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.203.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.204.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.204.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.204.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.205.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.205.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.205.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.206.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.206.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.206.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.207.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.207.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.207.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.208.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.208.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.208.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.209.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.209.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.209.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.210.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.210.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.210.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.211.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.211.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.211.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.212.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.212.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.212.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.213.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.213.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.213.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.214.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.214.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.214.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.215.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.215.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.215.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.216.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.216.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.216.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.217.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.217.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.217.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.218.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.218.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.218.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.219.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.219.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.219.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.220.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.220.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.220.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.221.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.221.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.221.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.222.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.222.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.222.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.223.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.223.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.223.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.224.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.224.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.224.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.225.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.225.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.225.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.226.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.226.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.226.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.227.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.227.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.227.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.228.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.228.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.228.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.229.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.229.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.229.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.230.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.230.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.230.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.231.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.231.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.231.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.232.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.232.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.232.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.233.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.233.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.233.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.234.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.234.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.234.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.235.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.235.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.235.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.236.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.236.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.236.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.237.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.237.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.237.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.238.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.238.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.238.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.239.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.239.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.239.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.240.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.240.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.240.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.241.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.241.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.241.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.242.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.242.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.242.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.243.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.243.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.243.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.244.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.244.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.244.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.245.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.245.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.245.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.246.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.246.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.246.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.247.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.247.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.247.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.248.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.248.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.248.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.249.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.249.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.249.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.250.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.250.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.250.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.251.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.251.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.251.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.252.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.252.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.252.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.253.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.253.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.253.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.254.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.254.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.254.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.255.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.255.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.mlp.experts.255.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.26.input_layernorm.weight": "model-00067-of-000163.safetensors", + "model.layers.26.post_attention_layernorm.weight": "model-00067-of-000163.safetensors", + "model.layers.27.self_attn.q_a_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.self_attn.q_a_layernorm.weight": "model-00067-of-000163.safetensors", + "model.layers.27.self_attn.q_b_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.self_attn.kv_a_proj_with_mqa.weight": "model-00067-of-000163.safetensors", + "model.layers.27.self_attn.kv_a_layernorm.weight": "model-00067-of-000163.safetensors", + "model.layers.27.self_attn.kv_b_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.self_attn.o_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.gate.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.gate.e_score_correction_bias": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.shared_experts.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.shared_experts.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.shared_experts.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.0.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.0.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.0.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.1.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.1.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.1.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.2.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.2.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.2.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.3.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.3.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.3.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.4.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.4.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.4.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.5.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.5.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.5.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.6.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.6.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.6.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.7.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.7.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.7.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.8.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.8.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.8.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.9.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.9.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.9.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.10.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.10.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.10.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.11.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.11.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.11.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.12.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.12.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.12.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.13.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.13.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.13.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.14.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.14.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.14.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.15.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.15.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.15.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.16.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.16.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.16.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.17.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.17.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.17.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.18.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.18.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.18.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.19.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.19.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.19.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.20.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.20.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.20.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.21.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.21.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.21.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.22.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.22.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.22.down_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.23.gate_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.23.up_proj.weight": "model-00067-of-000163.safetensors", + "model.layers.27.mlp.experts.23.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.24.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.24.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.24.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.25.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.25.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.25.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.26.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.26.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.26.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.27.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.27.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.27.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.28.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.28.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.28.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.29.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.29.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.29.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.30.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.30.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.30.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.31.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.31.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.31.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.32.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.32.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.32.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.33.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.33.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.33.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.34.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.34.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.34.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.35.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.35.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.35.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.36.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.36.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.36.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.37.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.37.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.37.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.38.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.38.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.38.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.39.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.39.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.39.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.40.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.40.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.40.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.41.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.41.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.41.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.42.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.42.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.42.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.43.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.43.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.43.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.44.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.44.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.44.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.45.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.45.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.45.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.46.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.46.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.46.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.47.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.47.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.47.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.48.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.48.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.48.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.49.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.49.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.49.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.50.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.50.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.50.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.51.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.51.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.51.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.52.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.52.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.52.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.53.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.53.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.53.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.54.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.54.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.54.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.55.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.55.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.55.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.56.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.56.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.56.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.57.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.57.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.57.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.58.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.58.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.58.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.59.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.59.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.59.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.60.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.60.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.60.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.61.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.61.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.61.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.62.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.62.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.62.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.63.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.63.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.63.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.64.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.64.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.64.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.65.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.65.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.65.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.66.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.66.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.66.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.67.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.67.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.67.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.68.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.68.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.68.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.69.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.69.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.69.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.70.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.70.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.70.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.71.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.71.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.71.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.72.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.72.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.72.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.73.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.73.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.73.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.74.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.74.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.74.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.75.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.75.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.75.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.76.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.76.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.76.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.77.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.77.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.77.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.78.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.78.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.78.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.79.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.79.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.79.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.80.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.80.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.80.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.81.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.81.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.81.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.82.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.82.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.82.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.83.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.83.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.83.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.84.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.84.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.84.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.85.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.85.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.85.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.86.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.86.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.86.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.87.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.87.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.87.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.88.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.88.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.88.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.89.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.89.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.89.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.90.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.90.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.90.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.91.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.91.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.91.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.92.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.92.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.92.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.93.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.93.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.93.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.94.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.94.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.94.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.95.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.95.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.95.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.96.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.96.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.96.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.97.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.97.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.97.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.98.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.98.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.98.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.99.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.99.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.99.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.100.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.100.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.100.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.101.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.101.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.101.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.102.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.102.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.102.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.103.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.103.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.103.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.104.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.104.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.104.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.105.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.105.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.105.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.106.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.106.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.106.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.107.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.107.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.107.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.108.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.108.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.108.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.109.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.109.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.109.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.110.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.110.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.110.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.111.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.111.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.111.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.112.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.112.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.112.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.113.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.113.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.113.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.114.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.114.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.114.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.115.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.115.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.115.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.116.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.116.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.116.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.117.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.117.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.117.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.118.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.118.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.118.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.119.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.119.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.119.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.120.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.120.up_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.120.down_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.121.gate_proj.weight": "model-00068-of-000163.safetensors", + "model.layers.27.mlp.experts.121.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.121.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.122.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.122.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.122.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.123.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.123.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.123.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.124.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.124.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.124.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.125.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.125.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.125.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.126.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.126.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.126.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.127.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.127.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.127.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.128.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.128.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.128.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.129.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.129.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.129.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.130.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.130.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.130.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.131.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.131.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.131.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.132.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.132.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.132.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.133.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.133.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.133.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.134.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.134.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.134.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.135.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.135.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.135.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.136.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.136.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.136.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.137.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.137.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.137.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.138.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.138.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.138.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.139.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.139.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.139.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.140.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.140.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.140.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.141.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.141.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.141.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.142.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.142.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.142.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.143.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.143.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.143.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.144.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.144.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.144.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.145.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.145.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.145.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.146.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.146.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.146.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.147.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.147.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.147.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.148.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.148.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.148.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.149.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.149.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.149.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.150.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.150.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.150.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.151.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.151.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.151.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.152.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.152.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.152.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.153.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.153.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.153.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.154.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.154.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.154.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.155.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.155.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.155.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.156.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.156.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.156.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.157.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.157.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.157.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.158.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.158.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.158.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.159.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.159.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.159.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.160.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.160.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.160.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.161.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.161.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.161.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.162.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.162.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.162.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.163.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.163.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.163.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.164.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.164.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.164.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.165.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.165.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.165.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.166.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.166.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.166.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.167.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.167.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.167.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.168.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.168.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.168.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.169.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.169.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.169.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.170.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.170.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.170.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.171.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.171.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.171.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.172.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.172.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.172.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.173.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.173.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.173.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.174.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.174.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.174.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.175.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.175.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.175.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.176.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.176.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.176.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.177.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.177.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.177.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.178.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.178.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.178.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.179.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.179.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.179.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.180.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.180.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.180.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.181.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.181.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.181.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.182.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.182.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.182.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.183.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.183.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.183.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.184.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.184.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.184.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.185.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.185.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.185.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.186.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.186.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.186.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.187.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.187.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.187.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.188.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.188.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.188.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.189.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.189.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.189.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.190.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.190.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.190.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.191.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.191.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.191.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.192.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.192.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.192.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.193.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.193.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.193.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.194.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.194.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.194.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.195.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.195.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.195.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.196.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.196.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.196.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.197.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.197.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.197.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.198.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.198.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.198.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.199.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.199.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.199.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.200.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.200.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.200.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.201.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.201.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.201.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.202.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.202.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.202.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.203.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.203.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.203.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.204.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.204.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.204.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.205.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.205.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.205.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.206.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.206.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.206.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.207.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.207.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.207.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.208.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.208.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.208.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.209.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.209.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.209.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.210.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.210.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.210.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.211.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.211.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.211.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.212.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.212.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.212.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.213.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.213.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.213.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.214.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.214.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.214.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.215.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.215.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.215.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.216.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.216.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.216.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.217.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.217.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.217.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.218.gate_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.218.up_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.218.down_proj.weight": "model-00069-of-000163.safetensors", + "model.layers.27.mlp.experts.219.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.219.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.219.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.220.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.220.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.220.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.221.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.221.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.221.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.222.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.222.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.222.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.223.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.223.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.223.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.224.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.224.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.224.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.225.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.225.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.225.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.226.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.226.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.226.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.227.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.227.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.227.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.228.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.228.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.228.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.229.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.229.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.229.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.230.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.230.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.230.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.231.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.231.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.231.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.232.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.232.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.232.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.233.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.233.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.233.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.234.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.234.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.234.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.235.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.235.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.235.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.236.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.236.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.236.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.237.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.237.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.237.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.238.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.238.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.238.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.239.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.239.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.239.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.240.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.240.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.240.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.241.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.241.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.241.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.242.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.242.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.242.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.243.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.243.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.243.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.244.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.244.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.244.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.245.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.245.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.245.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.246.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.246.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.246.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.247.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.247.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.247.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.248.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.248.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.248.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.249.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.249.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.249.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.250.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.250.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.250.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.251.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.251.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.251.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.252.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.252.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.252.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.253.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.253.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.253.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.254.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.254.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.254.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.255.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.255.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.mlp.experts.255.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.27.input_layernorm.weight": "model-00070-of-000163.safetensors", + "model.layers.27.post_attention_layernorm.weight": "model-00070-of-000163.safetensors", + "model.layers.28.self_attn.q_a_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.self_attn.q_a_layernorm.weight": "model-00070-of-000163.safetensors", + "model.layers.28.self_attn.q_b_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.self_attn.kv_a_proj_with_mqa.weight": "model-00070-of-000163.safetensors", + "model.layers.28.self_attn.kv_a_layernorm.weight": "model-00070-of-000163.safetensors", + "model.layers.28.self_attn.kv_b_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.self_attn.o_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.gate.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.gate.e_score_correction_bias": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.shared_experts.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.shared_experts.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.shared_experts.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.0.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.0.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.0.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.1.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.1.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.1.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.2.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.2.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.2.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.3.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.3.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.3.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.4.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.4.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.4.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.5.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.5.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.5.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.6.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.6.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.6.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.7.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.7.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.7.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.8.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.8.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.8.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.9.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.9.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.9.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.10.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.10.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.10.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.11.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.11.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.11.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.12.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.12.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.12.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.13.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.13.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.13.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.14.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.14.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.14.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.15.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.15.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.15.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.16.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.16.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.16.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.17.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.17.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.17.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.18.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.18.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.18.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.19.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.19.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.19.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.20.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.20.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.20.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.21.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.21.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.21.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.22.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.22.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.22.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.23.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.23.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.23.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.24.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.24.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.24.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.25.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.25.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.25.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.26.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.26.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.26.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.27.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.27.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.27.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.28.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.28.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.28.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.29.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.29.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.29.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.30.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.30.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.30.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.31.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.31.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.31.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.32.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.32.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.32.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.33.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.33.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.33.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.34.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.34.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.34.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.35.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.35.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.35.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.36.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.36.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.36.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.37.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.37.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.37.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.38.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.38.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.38.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.39.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.39.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.39.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.40.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.40.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.40.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.41.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.41.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.41.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.42.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.42.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.42.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.43.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.43.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.43.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.44.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.44.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.44.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.45.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.45.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.45.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.46.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.46.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.46.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.47.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.47.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.47.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.48.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.48.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.48.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.49.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.49.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.49.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.50.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.50.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.50.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.51.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.51.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.51.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.52.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.52.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.52.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.53.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.53.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.53.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.54.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.54.up_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.54.down_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.55.gate_proj.weight": "model-00070-of-000163.safetensors", + "model.layers.28.mlp.experts.55.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.55.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.56.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.56.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.56.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.57.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.57.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.57.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.58.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.58.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.58.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.59.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.59.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.59.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.60.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.60.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.60.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.61.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.61.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.61.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.62.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.62.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.62.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.63.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.63.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.63.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.64.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.64.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.64.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.65.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.65.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.65.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.66.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.66.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.66.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.67.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.67.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.67.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.68.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.68.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.68.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.69.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.69.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.69.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.70.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.70.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.70.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.71.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.71.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.71.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.72.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.72.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.72.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.73.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.73.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.73.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.74.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.74.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.74.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.75.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.75.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.75.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.76.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.76.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.76.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.77.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.77.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.77.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.78.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.78.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.78.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.79.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.79.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.79.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.80.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.80.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.80.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.81.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.81.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.81.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.82.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.82.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.82.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.83.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.83.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.83.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.84.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.84.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.84.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.85.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.85.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.85.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.86.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.86.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.86.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.87.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.87.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.87.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.88.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.88.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.88.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.89.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.89.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.89.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.90.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.90.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.90.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.91.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.91.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.91.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.92.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.92.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.92.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.93.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.93.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.93.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.94.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.94.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.94.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.95.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.95.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.95.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.96.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.96.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.96.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.97.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.97.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.97.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.98.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.98.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.98.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.99.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.99.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.99.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.100.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.100.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.100.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.101.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.101.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.101.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.102.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.102.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.102.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.103.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.103.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.103.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.104.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.104.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.104.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.105.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.105.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.105.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.106.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.106.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.106.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.107.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.107.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.107.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.108.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.108.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.108.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.109.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.109.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.109.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.110.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.110.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.110.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.111.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.111.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.111.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.112.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.112.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.112.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.113.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.113.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.113.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.114.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.114.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.114.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.115.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.115.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.115.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.116.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.116.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.116.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.117.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.117.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.117.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.118.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.118.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.118.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.119.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.119.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.119.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.120.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.120.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.120.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.121.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.121.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.121.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.122.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.122.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.122.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.123.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.123.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.123.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.124.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.124.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.124.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.125.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.125.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.125.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.126.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.126.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.126.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.127.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.127.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.127.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.128.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.128.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.128.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.129.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.129.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.129.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.130.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.130.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.130.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.131.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.131.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.131.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.132.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.132.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.132.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.133.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.133.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.133.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.134.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.134.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.134.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.135.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.135.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.135.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.136.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.136.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.136.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.137.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.137.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.137.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.138.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.138.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.138.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.139.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.139.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.139.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.140.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.140.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.140.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.141.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.141.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.141.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.142.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.142.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.142.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.143.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.143.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.143.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.144.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.144.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.144.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.145.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.145.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.145.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.146.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.146.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.146.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.147.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.147.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.147.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.148.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.148.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.148.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.149.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.149.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.149.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.150.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.150.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.150.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.151.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.151.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.151.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.152.gate_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.152.up_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.152.down_proj.weight": "model-00071-of-000163.safetensors", + "model.layers.28.mlp.experts.153.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.153.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.153.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.154.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.154.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.154.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.155.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.155.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.155.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.156.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.156.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.156.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.157.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.157.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.157.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.158.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.158.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.158.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.159.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.159.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.159.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.160.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.160.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.160.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.161.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.161.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.161.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.162.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.162.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.162.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.163.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.163.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.163.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.164.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.164.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.164.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.165.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.165.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.165.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.166.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.166.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.166.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.167.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.167.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.167.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.168.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.168.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.168.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.169.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.169.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.169.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.170.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.170.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.170.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.171.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.171.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.171.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.172.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.172.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.172.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.173.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.173.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.173.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.174.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.174.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.174.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.175.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.175.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.175.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.176.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.176.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.176.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.177.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.177.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.177.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.178.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.178.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.178.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.179.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.179.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.179.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.180.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.180.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.180.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.181.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.181.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.181.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.182.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.182.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.182.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.183.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.183.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.183.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.184.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.184.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.184.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.185.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.185.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.185.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.186.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.186.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.186.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.187.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.187.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.187.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.188.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.188.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.188.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.189.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.189.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.189.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.190.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.190.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.190.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.191.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.191.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.191.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.192.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.192.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.192.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.193.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.193.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.193.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.194.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.194.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.194.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.195.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.195.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.195.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.196.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.196.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.196.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.197.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.197.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.197.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.198.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.198.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.198.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.199.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.199.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.199.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.200.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.200.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.200.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.201.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.201.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.201.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.202.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.202.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.202.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.203.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.203.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.203.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.204.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.204.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.204.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.205.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.205.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.205.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.206.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.206.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.206.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.207.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.207.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.207.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.208.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.208.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.208.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.209.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.209.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.209.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.210.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.210.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.210.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.211.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.211.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.211.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.212.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.212.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.212.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.213.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.213.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.213.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.214.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.214.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.214.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.215.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.215.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.215.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.216.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.216.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.216.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.217.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.217.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.217.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.218.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.218.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.218.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.219.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.219.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.219.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.220.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.220.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.220.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.221.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.221.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.221.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.222.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.222.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.222.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.223.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.223.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.223.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.224.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.224.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.224.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.225.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.225.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.225.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.226.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.226.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.226.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.227.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.227.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.227.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.228.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.228.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.228.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.229.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.229.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.229.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.230.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.230.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.230.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.231.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.231.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.231.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.232.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.232.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.232.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.233.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.233.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.233.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.234.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.234.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.234.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.235.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.235.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.235.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.236.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.236.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.236.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.237.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.237.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.237.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.238.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.238.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.238.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.239.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.239.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.239.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.240.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.240.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.240.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.241.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.241.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.241.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.242.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.242.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.242.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.243.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.243.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.243.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.244.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.244.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.244.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.245.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.245.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.245.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.246.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.246.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.246.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.247.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.247.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.247.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.248.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.248.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.248.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.249.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.249.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.249.down_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.250.gate_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.250.up_proj.weight": "model-00072-of-000163.safetensors", + "model.layers.28.mlp.experts.250.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.28.mlp.experts.251.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.28.mlp.experts.251.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.28.mlp.experts.251.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.28.mlp.experts.252.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.28.mlp.experts.252.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.28.mlp.experts.252.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.28.mlp.experts.253.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.28.mlp.experts.253.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.28.mlp.experts.253.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.28.mlp.experts.254.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.28.mlp.experts.254.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.28.mlp.experts.254.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.28.mlp.experts.255.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.28.mlp.experts.255.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.28.mlp.experts.255.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.28.input_layernorm.weight": "model-00073-of-000163.safetensors", + "model.layers.28.post_attention_layernorm.weight": "model-00073-of-000163.safetensors", + "model.layers.29.self_attn.q_a_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.self_attn.q_a_layernorm.weight": "model-00073-of-000163.safetensors", + "model.layers.29.self_attn.q_b_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.self_attn.kv_a_proj_with_mqa.weight": "model-00073-of-000163.safetensors", + "model.layers.29.self_attn.kv_a_layernorm.weight": "model-00073-of-000163.safetensors", + "model.layers.29.self_attn.kv_b_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.self_attn.o_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.gate.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.gate.e_score_correction_bias": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.shared_experts.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.shared_experts.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.shared_experts.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.0.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.0.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.0.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.1.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.1.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.1.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.2.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.2.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.2.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.3.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.3.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.3.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.4.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.4.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.4.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.5.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.5.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.5.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.6.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.6.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.6.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.7.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.7.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.7.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.8.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.8.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.8.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.9.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.9.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.9.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.10.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.10.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.10.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.11.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.11.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.11.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.12.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.12.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.12.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.13.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.13.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.13.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.14.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.14.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.14.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.15.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.15.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.15.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.16.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.16.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.16.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.17.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.17.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.17.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.18.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.18.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.18.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.19.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.19.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.19.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.20.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.20.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.20.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.21.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.21.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.21.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.22.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.22.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.22.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.23.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.23.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.23.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.24.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.24.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.24.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.25.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.25.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.25.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.26.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.26.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.26.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.27.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.27.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.27.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.28.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.28.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.28.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.29.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.29.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.29.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.30.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.30.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.30.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.31.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.31.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.31.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.32.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.32.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.32.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.33.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.33.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.33.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.34.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.34.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.34.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.35.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.35.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.35.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.36.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.36.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.36.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.37.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.37.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.37.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.38.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.38.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.38.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.39.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.39.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.39.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.40.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.40.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.40.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.41.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.41.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.41.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.42.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.42.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.42.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.43.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.43.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.43.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.44.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.44.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.44.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.45.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.45.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.45.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.46.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.46.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.46.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.47.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.47.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.47.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.48.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.48.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.48.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.49.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.49.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.49.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.50.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.50.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.50.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.51.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.51.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.51.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.52.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.52.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.52.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.53.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.53.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.53.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.54.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.54.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.54.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.55.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.55.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.55.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.56.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.56.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.56.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.57.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.57.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.57.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.58.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.58.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.58.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.59.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.59.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.59.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.60.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.60.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.60.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.61.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.61.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.61.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.62.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.62.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.62.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.63.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.63.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.63.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.64.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.64.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.64.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.65.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.65.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.65.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.66.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.66.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.66.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.67.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.67.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.67.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.68.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.68.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.68.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.69.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.69.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.69.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.70.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.70.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.70.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.71.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.71.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.71.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.72.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.72.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.72.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.73.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.73.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.73.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.74.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.74.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.74.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.75.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.75.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.75.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.76.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.76.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.76.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.77.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.77.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.77.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.78.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.78.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.78.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.79.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.79.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.79.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.80.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.80.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.80.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.81.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.81.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.81.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.82.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.82.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.82.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.83.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.83.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.83.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.84.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.84.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.84.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.85.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.85.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.85.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.86.gate_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.86.up_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.86.down_proj.weight": "model-00073-of-000163.safetensors", + "model.layers.29.mlp.experts.87.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.87.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.87.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.88.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.88.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.88.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.89.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.89.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.89.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.90.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.90.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.90.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.91.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.91.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.91.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.92.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.92.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.92.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.93.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.93.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.93.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.94.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.94.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.94.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.95.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.95.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.95.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.96.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.96.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.96.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.97.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.97.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.97.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.98.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.98.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.98.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.99.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.99.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.99.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.100.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.100.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.100.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.101.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.101.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.101.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.102.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.102.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.102.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.103.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.103.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.103.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.104.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.104.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.104.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.105.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.105.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.105.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.106.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.106.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.106.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.107.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.107.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.107.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.108.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.108.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.108.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.109.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.109.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.109.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.110.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.110.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.110.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.111.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.111.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.111.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.112.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.112.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.112.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.113.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.113.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.113.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.114.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.114.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.114.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.115.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.115.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.115.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.116.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.116.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.116.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.117.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.117.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.117.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.118.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.118.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.118.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.119.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.119.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.119.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.120.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.120.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.120.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.121.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.121.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.121.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.122.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.122.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.122.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.123.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.123.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.123.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.124.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.124.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.124.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.125.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.125.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.125.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.126.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.126.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.126.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.127.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.127.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.127.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.128.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.128.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.128.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.129.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.129.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.129.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.130.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.130.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.130.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.131.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.131.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.131.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.132.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.132.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.132.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.133.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.133.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.133.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.134.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.134.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.134.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.135.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.135.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.135.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.136.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.136.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.136.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.137.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.137.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.137.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.138.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.138.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.138.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.139.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.139.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.139.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.140.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.140.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.140.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.141.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.141.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.141.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.142.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.142.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.142.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.143.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.143.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.143.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.144.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.144.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.144.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.145.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.145.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.145.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.146.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.146.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.146.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.147.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.147.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.147.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.148.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.148.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.148.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.149.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.149.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.149.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.150.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.150.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.150.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.151.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.151.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.151.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.152.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.152.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.152.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.153.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.153.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.153.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.154.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.154.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.154.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.155.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.155.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.155.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.156.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.156.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.156.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.157.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.157.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.157.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.158.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.158.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.158.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.159.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.159.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.159.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.160.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.160.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.160.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.161.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.161.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.161.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.162.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.162.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.162.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.163.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.163.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.163.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.164.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.164.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.164.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.165.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.165.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.165.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.166.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.166.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.166.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.167.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.167.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.167.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.168.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.168.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.168.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.169.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.169.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.169.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.170.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.170.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.170.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.171.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.171.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.171.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.172.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.172.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.172.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.173.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.173.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.173.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.174.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.174.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.174.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.175.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.175.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.175.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.176.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.176.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.176.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.177.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.177.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.177.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.178.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.178.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.178.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.179.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.179.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.179.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.180.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.180.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.180.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.181.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.181.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.181.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.182.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.182.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.182.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.183.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.183.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.183.down_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.184.gate_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.184.up_proj.weight": "model-00074-of-000163.safetensors", + "model.layers.29.mlp.experts.184.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.185.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.185.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.185.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.186.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.186.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.186.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.187.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.187.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.187.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.188.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.188.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.188.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.189.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.189.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.189.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.190.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.190.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.190.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.191.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.191.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.191.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.192.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.192.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.192.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.193.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.193.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.193.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.194.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.194.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.194.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.195.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.195.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.195.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.196.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.196.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.196.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.197.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.197.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.197.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.198.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.198.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.198.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.199.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.199.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.199.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.200.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.200.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.200.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.201.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.201.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.201.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.202.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.202.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.202.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.203.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.203.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.203.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.204.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.204.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.204.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.205.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.205.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.205.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.206.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.206.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.206.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.207.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.207.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.207.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.208.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.208.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.208.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.209.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.209.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.209.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.210.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.210.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.210.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.211.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.211.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.211.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.212.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.212.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.212.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.213.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.213.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.213.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.214.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.214.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.214.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.215.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.215.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.215.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.216.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.216.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.216.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.217.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.217.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.217.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.218.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.218.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.218.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.219.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.219.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.219.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.220.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.220.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.220.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.221.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.221.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.221.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.222.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.222.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.222.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.223.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.223.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.223.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.224.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.224.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.224.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.225.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.225.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.225.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.226.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.226.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.226.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.227.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.227.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.227.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.228.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.228.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.228.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.229.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.229.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.229.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.230.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.230.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.230.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.231.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.231.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.231.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.232.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.232.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.232.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.233.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.233.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.233.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.234.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.234.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.234.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.235.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.235.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.235.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.236.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.236.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.236.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.237.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.237.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.237.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.238.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.238.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.238.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.239.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.239.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.239.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.240.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.240.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.240.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.241.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.241.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.241.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.242.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.242.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.242.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.243.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.243.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.243.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.244.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.244.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.244.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.245.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.245.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.245.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.246.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.246.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.246.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.247.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.247.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.247.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.248.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.248.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.248.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.249.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.249.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.249.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.250.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.250.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.250.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.251.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.251.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.251.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.252.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.252.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.252.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.253.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.253.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.253.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.254.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.254.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.254.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.255.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.255.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.mlp.experts.255.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.29.input_layernorm.weight": "model-00075-of-000163.safetensors", + "model.layers.29.post_attention_layernorm.weight": "model-00075-of-000163.safetensors", + "model.layers.30.self_attn.q_a_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.self_attn.q_a_layernorm.weight": "model-00075-of-000163.safetensors", + "model.layers.30.self_attn.q_b_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.self_attn.kv_a_proj_with_mqa.weight": "model-00075-of-000163.safetensors", + "model.layers.30.self_attn.kv_a_layernorm.weight": "model-00075-of-000163.safetensors", + "model.layers.30.self_attn.kv_b_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.self_attn.o_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.gate.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.gate.e_score_correction_bias": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.shared_experts.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.shared_experts.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.shared_experts.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.0.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.0.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.0.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.1.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.1.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.1.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.2.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.2.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.2.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.3.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.3.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.3.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.4.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.4.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.4.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.5.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.5.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.5.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.6.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.6.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.6.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.7.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.7.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.7.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.8.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.8.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.8.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.9.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.9.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.9.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.10.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.10.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.10.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.11.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.11.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.11.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.12.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.12.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.12.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.13.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.13.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.13.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.14.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.14.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.14.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.15.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.15.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.15.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.16.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.16.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.16.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.17.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.17.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.17.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.18.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.18.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.18.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.19.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.19.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.19.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.20.gate_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.20.up_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.20.down_proj.weight": "model-00075-of-000163.safetensors", + "model.layers.30.mlp.experts.21.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.21.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.21.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.22.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.22.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.22.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.23.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.23.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.23.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.24.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.24.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.24.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.25.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.25.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.25.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.26.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.26.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.26.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.27.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.27.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.27.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.28.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.28.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.28.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.29.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.29.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.29.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.30.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.30.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.30.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.31.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.31.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.31.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.32.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.32.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.32.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.33.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.33.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.33.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.34.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.34.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.34.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.35.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.35.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.35.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.36.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.36.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.36.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.37.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.37.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.37.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.38.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.38.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.38.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.39.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.39.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.39.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.40.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.40.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.40.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.41.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.41.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.41.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.42.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.42.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.42.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.43.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.43.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.43.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.44.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.44.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.44.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.45.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.45.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.45.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.46.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.46.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.46.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.47.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.47.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.47.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.48.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.48.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.48.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.49.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.49.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.49.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.50.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.50.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.50.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.51.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.51.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.51.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.52.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.52.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.52.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.53.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.53.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.53.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.54.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.54.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.54.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.55.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.55.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.55.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.56.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.56.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.56.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.57.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.57.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.57.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.58.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.58.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.58.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.59.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.59.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.59.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.60.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.60.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.60.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.61.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.61.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.61.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.62.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.62.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.62.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.63.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.63.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.63.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.64.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.64.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.64.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.65.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.65.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.65.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.66.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.66.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.66.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.67.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.67.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.67.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.68.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.68.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.68.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.69.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.69.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.69.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.70.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.70.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.70.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.71.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.71.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.71.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.72.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.72.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.72.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.73.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.73.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.73.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.74.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.74.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.74.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.75.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.75.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.75.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.76.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.76.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.76.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.77.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.77.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.77.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.78.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.78.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.78.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.79.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.79.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.79.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.80.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.80.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.80.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.81.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.81.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.81.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.82.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.82.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.82.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.83.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.83.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.83.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.84.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.84.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.84.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.85.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.85.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.85.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.86.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.86.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.86.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.87.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.87.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.87.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.88.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.88.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.88.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.89.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.89.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.89.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.90.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.90.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.90.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.91.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.91.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.91.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.92.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.92.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.92.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.93.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.93.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.93.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.94.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.94.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.94.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.95.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.95.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.95.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.96.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.96.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.96.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.97.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.97.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.97.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.98.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.98.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.98.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.99.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.99.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.99.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.100.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.100.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.100.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.101.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.101.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.101.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.102.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.102.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.102.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.103.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.103.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.103.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.104.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.104.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.104.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.105.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.105.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.105.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.106.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.106.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.106.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.107.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.107.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.107.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.108.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.108.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.108.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.109.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.109.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.109.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.110.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.110.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.110.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.111.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.111.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.111.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.112.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.112.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.112.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.113.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.113.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.113.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.114.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.114.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.114.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.115.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.115.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.115.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.116.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.116.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.116.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.117.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.117.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.117.down_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.118.gate_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.118.up_proj.weight": "model-00076-of-000163.safetensors", + "model.layers.30.mlp.experts.118.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.119.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.119.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.119.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.120.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.120.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.120.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.121.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.121.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.121.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.122.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.122.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.122.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.123.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.123.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.123.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.124.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.124.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.124.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.125.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.125.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.125.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.126.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.126.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.126.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.127.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.127.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.127.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.128.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.128.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.128.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.129.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.129.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.129.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.130.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.130.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.130.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.131.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.131.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.131.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.132.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.132.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.132.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.133.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.133.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.133.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.134.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.134.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.134.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.135.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.135.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.135.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.136.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.136.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.136.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.137.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.137.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.137.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.138.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.138.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.138.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.139.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.139.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.139.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.140.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.140.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.140.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.141.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.141.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.141.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.142.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.142.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.142.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.143.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.143.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.143.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.144.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.144.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.144.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.145.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.145.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.145.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.146.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.146.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.146.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.147.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.147.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.147.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.148.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.148.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.148.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.149.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.149.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.149.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.150.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.150.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.150.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.151.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.151.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.151.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.152.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.152.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.152.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.153.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.153.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.153.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.154.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.154.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.154.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.155.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.155.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.155.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.156.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.156.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.156.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.157.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.157.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.157.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.158.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.158.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.158.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.159.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.159.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.159.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.160.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.160.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.160.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.161.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.161.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.161.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.162.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.162.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.162.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.163.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.163.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.163.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.164.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.164.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.164.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.165.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.165.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.165.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.166.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.166.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.166.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.167.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.167.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.167.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.168.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.168.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.168.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.169.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.169.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.169.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.170.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.170.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.170.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.171.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.171.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.171.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.172.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.172.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.172.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.173.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.173.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.173.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.174.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.174.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.174.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.175.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.175.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.175.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.176.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.176.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.176.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.177.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.177.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.177.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.178.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.178.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.178.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.179.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.179.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.179.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.180.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.180.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.180.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.181.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.181.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.181.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.182.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.182.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.182.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.183.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.183.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.183.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.184.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.184.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.184.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.185.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.185.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.185.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.186.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.186.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.186.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.187.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.187.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.187.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.188.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.188.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.188.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.189.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.189.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.189.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.190.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.190.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.190.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.191.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.191.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.191.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.192.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.192.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.192.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.193.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.193.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.193.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.194.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.194.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.194.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.195.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.195.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.195.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.196.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.196.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.196.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.197.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.197.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.197.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.198.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.198.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.198.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.199.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.199.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.199.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.200.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.200.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.200.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.201.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.201.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.201.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.202.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.202.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.202.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.203.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.203.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.203.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.204.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.204.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.204.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.205.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.205.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.205.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.206.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.206.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.206.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.207.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.207.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.207.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.208.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.208.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.208.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.209.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.209.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.209.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.210.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.210.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.210.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.211.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.211.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.211.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.212.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.212.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.212.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.213.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.213.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.213.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.214.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.214.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.214.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.215.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.215.up_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.215.down_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.216.gate_proj.weight": "model-00077-of-000163.safetensors", + "model.layers.30.mlp.experts.216.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.216.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.217.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.217.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.217.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.218.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.218.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.218.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.219.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.219.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.219.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.220.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.220.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.220.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.221.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.221.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.221.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.222.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.222.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.222.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.223.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.223.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.223.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.224.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.224.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.224.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.225.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.225.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.225.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.226.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.226.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.226.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.227.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.227.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.227.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.228.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.228.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.228.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.229.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.229.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.229.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.230.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.230.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.230.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.231.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.231.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.231.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.232.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.232.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.232.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.233.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.233.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.233.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.234.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.234.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.234.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.235.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.235.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.235.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.236.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.236.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.236.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.237.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.237.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.237.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.238.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.238.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.238.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.239.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.239.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.239.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.240.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.240.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.240.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.241.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.241.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.241.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.242.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.242.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.242.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.243.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.243.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.243.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.244.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.244.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.244.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.245.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.245.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.245.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.246.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.246.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.246.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.247.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.247.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.247.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.248.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.248.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.248.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.249.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.249.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.249.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.250.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.250.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.250.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.251.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.251.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.251.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.252.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.252.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.252.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.253.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.253.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.253.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.254.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.254.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.254.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.255.gate_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.255.up_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.mlp.experts.255.down_proj.weight": "model-00078-of-000163.safetensors", + "model.layers.30.input_layernorm.weight": "model-00078-of-000163.safetensors", + "model.layers.30.post_attention_layernorm.weight": "model-00078-of-000163.safetensors", + "model.layers.31.self_attn.q_a_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.self_attn.q_a_layernorm.weight": "model-00079-of-000163.safetensors", + "model.layers.31.self_attn.q_b_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.self_attn.kv_a_proj_with_mqa.weight": "model-00079-of-000163.safetensors", + "model.layers.31.self_attn.kv_a_layernorm.weight": "model-00079-of-000163.safetensors", + "model.layers.31.self_attn.kv_b_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.self_attn.o_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.gate.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.gate.e_score_correction_bias": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.shared_experts.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.shared_experts.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.shared_experts.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.0.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.0.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.0.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.1.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.1.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.1.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.2.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.2.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.2.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.3.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.3.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.3.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.4.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.4.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.4.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.5.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.5.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.5.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.6.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.6.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.6.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.7.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.7.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.7.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.8.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.8.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.8.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.9.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.9.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.9.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.10.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.10.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.10.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.11.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.11.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.11.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.12.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.12.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.12.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.13.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.13.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.13.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.14.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.14.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.14.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.15.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.15.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.15.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.16.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.16.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.16.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.17.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.17.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.17.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.18.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.18.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.18.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.19.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.19.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.19.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.20.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.20.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.20.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.21.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.21.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.21.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.22.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.22.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.22.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.23.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.23.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.23.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.24.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.24.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.24.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.25.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.25.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.25.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.26.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.26.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.26.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.27.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.27.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.27.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.28.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.28.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.28.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.29.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.29.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.29.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.30.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.30.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.30.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.31.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.31.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.31.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.32.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.32.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.32.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.33.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.33.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.33.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.34.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.34.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.34.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.35.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.35.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.35.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.36.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.36.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.36.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.37.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.37.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.37.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.38.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.38.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.38.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.39.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.39.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.39.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.40.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.40.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.40.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.41.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.41.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.41.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.42.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.42.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.42.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.43.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.43.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.43.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.44.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.44.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.44.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.45.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.45.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.45.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.46.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.46.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.46.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.47.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.47.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.47.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.48.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.48.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.48.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.49.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.49.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.49.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.50.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.50.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.50.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.51.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.51.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.51.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.52.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.52.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.52.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.53.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.53.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.53.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.54.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.54.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.54.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.55.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.55.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.55.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.56.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.56.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.56.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.57.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.57.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.57.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.58.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.58.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.58.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.59.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.59.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.59.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.60.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.60.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.60.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.61.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.61.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.61.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.62.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.62.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.62.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.63.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.63.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.63.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.64.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.64.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.64.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.65.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.65.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.65.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.66.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.66.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.66.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.67.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.67.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.67.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.68.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.68.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.68.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.69.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.69.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.69.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.70.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.70.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.70.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.71.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.71.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.71.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.72.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.72.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.72.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.73.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.73.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.73.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.74.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.74.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.74.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.75.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.75.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.75.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.76.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.76.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.76.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.77.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.77.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.77.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.78.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.78.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.78.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.79.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.79.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.79.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.80.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.80.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.80.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.81.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.81.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.81.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.82.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.82.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.82.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.83.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.83.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.83.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.84.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.84.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.84.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.85.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.85.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.85.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.86.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.86.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.86.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.87.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.87.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.87.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.88.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.88.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.88.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.89.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.89.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.89.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.90.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.90.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.90.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.91.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.91.up_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.91.down_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.92.gate_proj.weight": "model-00079-of-000163.safetensors", + "model.layers.31.mlp.experts.92.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.92.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.93.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.93.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.93.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.94.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.94.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.94.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.95.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.95.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.95.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.96.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.96.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.96.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.97.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.97.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.97.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.98.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.98.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.98.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.99.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.99.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.99.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.100.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.100.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.100.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.101.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.101.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.101.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.102.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.102.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.102.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.103.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.103.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.103.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.104.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.104.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.104.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.105.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.105.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.105.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.106.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.106.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.106.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.107.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.107.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.107.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.108.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.108.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.108.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.109.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.109.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.109.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.110.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.110.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.110.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.111.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.111.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.111.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.112.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.112.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.112.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.113.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.113.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.113.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.114.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.114.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.114.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.115.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.115.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.115.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.116.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.116.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.116.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.117.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.117.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.117.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.118.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.118.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.118.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.119.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.119.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.119.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.120.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.120.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.120.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.121.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.121.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.121.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.122.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.122.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.122.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.123.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.123.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.123.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.124.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.124.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.124.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.125.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.125.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.125.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.126.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.126.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.126.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.127.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.127.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.127.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.128.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.128.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.128.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.129.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.129.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.129.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.130.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.130.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.130.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.131.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.131.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.131.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.132.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.132.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.132.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.133.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.133.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.133.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.134.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.134.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.134.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.135.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.135.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.135.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.136.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.136.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.136.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.137.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.137.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.137.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.138.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.138.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.138.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.139.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.139.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.139.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.140.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.140.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.140.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.141.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.141.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.141.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.142.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.142.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.142.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.143.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.143.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.143.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.144.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.144.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.144.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.145.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.145.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.145.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.146.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.146.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.146.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.147.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.147.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.147.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.148.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.148.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.148.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.149.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.149.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.149.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.150.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.150.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.150.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.151.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.151.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.151.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.152.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.152.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.152.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.153.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.153.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.153.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.154.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.154.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.154.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.155.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.155.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.155.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.156.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.156.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.156.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.157.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.157.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.157.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.158.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.158.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.158.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.159.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.159.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.159.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.160.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.160.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.160.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.161.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.161.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.161.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.162.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.162.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.162.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.163.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.163.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.163.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.164.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.164.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.164.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.165.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.165.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.165.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.166.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.166.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.166.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.167.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.167.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.167.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.168.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.168.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.168.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.169.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.169.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.169.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.170.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.170.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.170.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.171.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.171.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.171.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.172.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.172.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.172.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.173.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.173.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.173.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.174.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.174.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.174.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.175.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.175.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.175.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.176.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.176.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.176.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.177.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.177.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.177.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.178.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.178.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.178.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.179.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.179.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.179.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.180.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.180.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.180.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.181.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.181.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.181.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.182.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.182.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.182.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.183.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.183.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.183.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.184.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.184.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.184.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.185.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.185.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.185.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.186.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.186.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.186.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.187.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.187.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.187.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.188.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.188.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.188.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.189.gate_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.189.up_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.189.down_proj.weight": "model-00080-of-000163.safetensors", + "model.layers.31.mlp.experts.190.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.190.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.190.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.191.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.191.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.191.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.192.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.192.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.192.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.193.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.193.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.193.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.194.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.194.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.194.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.195.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.195.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.195.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.196.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.196.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.196.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.197.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.197.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.197.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.198.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.198.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.198.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.199.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.199.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.199.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.200.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.200.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.200.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.201.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.201.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.201.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.202.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.202.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.202.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.203.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.203.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.203.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.204.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.204.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.204.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.205.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.205.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.205.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.206.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.206.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.206.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.207.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.207.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.207.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.208.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.208.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.208.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.209.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.209.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.209.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.210.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.210.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.210.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.211.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.211.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.211.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.212.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.212.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.212.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.213.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.213.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.213.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.214.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.214.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.214.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.215.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.215.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.215.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.216.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.216.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.216.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.217.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.217.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.217.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.218.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.218.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.218.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.219.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.219.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.219.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.220.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.220.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.220.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.221.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.221.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.221.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.222.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.222.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.222.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.223.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.223.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.223.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.224.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.224.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.224.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.225.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.225.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.225.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.226.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.226.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.226.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.227.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.227.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.227.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.228.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.228.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.228.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.229.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.229.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.229.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.230.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.230.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.230.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.231.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.231.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.231.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.232.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.232.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.232.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.233.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.233.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.233.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.234.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.234.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.234.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.235.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.235.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.235.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.236.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.236.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.236.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.237.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.237.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.237.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.238.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.238.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.238.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.239.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.239.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.239.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.240.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.240.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.240.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.241.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.241.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.241.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.242.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.242.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.242.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.243.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.243.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.243.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.244.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.244.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.244.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.245.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.245.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.245.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.246.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.246.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.246.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.247.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.247.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.247.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.248.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.248.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.248.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.249.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.249.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.249.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.250.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.250.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.250.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.251.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.251.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.251.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.252.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.252.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.252.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.253.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.253.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.253.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.254.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.254.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.254.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.255.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.255.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.mlp.experts.255.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.31.input_layernorm.weight": "model-00081-of-000163.safetensors", + "model.layers.31.post_attention_layernorm.weight": "model-00081-of-000163.safetensors", + "model.layers.32.self_attn.q_a_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.self_attn.q_a_layernorm.weight": "model-00081-of-000163.safetensors", + "model.layers.32.self_attn.q_b_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.self_attn.kv_a_proj_with_mqa.weight": "model-00081-of-000163.safetensors", + "model.layers.32.self_attn.kv_a_layernorm.weight": "model-00081-of-000163.safetensors", + "model.layers.32.self_attn.kv_b_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.self_attn.o_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.gate.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.gate.e_score_correction_bias": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.shared_experts.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.shared_experts.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.shared_experts.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.0.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.0.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.0.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.1.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.1.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.1.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.2.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.2.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.2.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.3.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.3.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.3.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.4.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.4.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.4.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.5.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.5.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.5.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.6.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.6.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.6.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.7.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.7.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.7.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.8.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.8.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.8.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.9.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.9.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.9.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.10.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.10.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.10.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.11.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.11.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.11.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.12.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.12.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.12.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.13.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.13.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.13.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.14.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.14.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.14.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.15.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.15.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.15.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.16.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.16.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.16.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.17.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.17.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.17.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.18.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.18.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.18.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.19.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.19.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.19.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.20.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.20.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.20.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.21.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.21.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.21.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.22.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.22.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.22.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.23.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.23.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.23.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.24.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.24.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.24.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.25.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.25.up_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.25.down_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.26.gate_proj.weight": "model-00081-of-000163.safetensors", + "model.layers.32.mlp.experts.26.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.26.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.27.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.27.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.27.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.28.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.28.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.28.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.29.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.29.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.29.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.30.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.30.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.30.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.31.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.31.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.31.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.32.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.32.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.32.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.33.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.33.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.33.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.34.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.34.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.34.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.35.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.35.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.35.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.36.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.36.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.36.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.37.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.37.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.37.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.38.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.38.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.38.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.39.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.39.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.39.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.40.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.40.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.40.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.41.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.41.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.41.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.42.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.42.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.42.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.43.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.43.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.43.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.44.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.44.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.44.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.45.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.45.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.45.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.46.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.46.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.46.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.47.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.47.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.47.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.48.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.48.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.48.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.49.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.49.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.49.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.50.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.50.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.50.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.51.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.51.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.51.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.52.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.52.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.52.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.53.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.53.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.53.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.54.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.54.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.54.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.55.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.55.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.55.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.56.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.56.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.56.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.57.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.57.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.57.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.58.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.58.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.58.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.59.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.59.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.59.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.60.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.60.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.60.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.61.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.61.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.61.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.62.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.62.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.62.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.63.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.63.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.63.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.64.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.64.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.64.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.65.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.65.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.65.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.66.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.66.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.66.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.67.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.67.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.67.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.68.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.68.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.68.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.69.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.69.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.69.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.70.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.70.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.70.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.71.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.71.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.71.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.72.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.72.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.72.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.73.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.73.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.73.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.74.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.74.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.74.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.75.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.75.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.75.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.76.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.76.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.76.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.77.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.77.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.77.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.78.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.78.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.78.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.79.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.79.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.79.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.80.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.80.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.80.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.81.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.81.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.81.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.82.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.82.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.82.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.83.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.83.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.83.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.84.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.84.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.84.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.85.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.85.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.85.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.86.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.86.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.86.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.87.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.87.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.87.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.88.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.88.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.88.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.89.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.89.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.89.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.90.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.90.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.90.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.91.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.91.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.91.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.92.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.92.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.92.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.93.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.93.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.93.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.94.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.94.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.94.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.95.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.95.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.95.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.96.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.96.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.96.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.97.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.97.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.97.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.98.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.98.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.98.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.99.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.99.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.99.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.100.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.100.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.100.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.101.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.101.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.101.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.102.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.102.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.102.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.103.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.103.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.103.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.104.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.104.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.104.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.105.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.105.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.105.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.106.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.106.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.106.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.107.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.107.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.107.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.108.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.108.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.108.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.109.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.109.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.109.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.110.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.110.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.110.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.111.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.111.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.111.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.112.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.112.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.112.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.113.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.113.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.113.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.114.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.114.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.114.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.115.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.115.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.115.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.116.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.116.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.116.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.117.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.117.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.117.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.118.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.118.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.118.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.119.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.119.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.119.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.120.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.120.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.120.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.121.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.121.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.121.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.122.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.122.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.122.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.123.gate_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.123.up_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.123.down_proj.weight": "model-00082-of-000163.safetensors", + "model.layers.32.mlp.experts.124.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.124.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.124.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.125.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.125.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.125.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.126.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.126.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.126.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.127.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.127.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.127.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.128.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.128.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.128.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.129.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.129.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.129.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.130.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.130.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.130.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.131.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.131.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.131.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.132.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.132.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.132.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.133.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.133.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.133.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.134.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.134.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.134.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.135.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.135.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.135.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.136.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.136.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.136.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.137.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.137.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.137.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.138.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.138.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.138.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.139.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.139.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.139.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.140.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.140.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.140.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.141.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.141.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.141.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.142.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.142.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.142.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.143.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.143.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.143.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.144.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.144.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.144.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.145.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.145.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.145.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.146.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.146.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.146.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.147.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.147.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.147.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.148.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.148.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.148.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.149.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.149.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.149.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.150.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.150.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.150.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.151.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.151.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.151.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.152.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.152.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.152.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.153.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.153.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.153.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.154.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.154.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.154.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.155.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.155.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.155.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.156.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.156.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.156.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.157.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.157.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.157.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.158.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.158.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.158.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.159.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.159.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.159.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.160.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.160.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.160.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.161.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.161.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.161.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.162.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.162.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.162.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.163.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.163.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.163.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.164.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.164.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.164.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.165.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.165.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.165.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.166.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.166.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.166.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.167.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.167.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.167.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.168.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.168.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.168.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.169.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.169.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.169.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.170.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.170.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.170.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.171.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.171.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.171.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.172.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.172.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.172.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.173.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.173.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.173.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.174.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.174.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.174.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.175.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.175.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.175.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.176.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.176.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.176.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.177.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.177.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.177.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.178.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.178.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.178.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.179.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.179.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.179.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.180.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.180.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.180.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.181.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.181.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.181.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.182.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.182.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.182.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.183.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.183.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.183.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.184.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.184.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.184.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.185.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.185.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.185.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.186.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.186.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.186.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.187.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.187.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.187.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.188.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.188.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.188.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.189.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.189.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.189.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.190.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.190.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.190.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.191.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.191.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.191.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.192.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.192.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.192.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.193.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.193.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.193.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.194.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.194.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.194.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.195.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.195.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.195.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.196.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.196.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.196.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.197.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.197.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.197.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.198.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.198.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.198.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.199.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.199.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.199.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.200.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.200.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.200.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.201.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.201.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.201.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.202.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.202.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.202.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.203.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.203.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.203.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.204.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.204.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.204.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.205.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.205.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.205.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.206.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.206.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.206.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.207.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.207.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.207.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.208.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.208.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.208.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.209.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.209.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.209.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.210.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.210.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.210.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.211.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.211.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.211.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.212.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.212.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.212.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.213.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.213.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.213.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.214.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.214.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.214.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.215.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.215.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.215.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.216.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.216.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.216.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.217.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.217.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.217.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.218.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.218.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.218.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.219.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.219.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.219.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.220.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.220.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.220.down_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.221.gate_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.221.up_proj.weight": "model-00083-of-000163.safetensors", + "model.layers.32.mlp.experts.221.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.222.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.222.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.222.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.223.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.223.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.223.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.224.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.224.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.224.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.225.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.225.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.225.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.226.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.226.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.226.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.227.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.227.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.227.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.228.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.228.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.228.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.229.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.229.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.229.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.230.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.230.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.230.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.231.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.231.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.231.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.232.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.232.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.232.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.233.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.233.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.233.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.234.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.234.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.234.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.235.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.235.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.235.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.236.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.236.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.236.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.237.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.237.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.237.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.238.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.238.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.238.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.239.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.239.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.239.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.240.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.240.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.240.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.241.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.241.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.241.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.242.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.242.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.242.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.243.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.243.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.243.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.244.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.244.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.244.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.245.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.245.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.245.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.246.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.246.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.246.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.247.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.247.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.247.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.248.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.248.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.248.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.249.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.249.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.249.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.250.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.250.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.250.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.251.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.251.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.251.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.252.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.252.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.252.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.253.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.253.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.253.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.254.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.254.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.254.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.255.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.255.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.mlp.experts.255.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.32.input_layernorm.weight": "model-00084-of-000163.safetensors", + "model.layers.32.post_attention_layernorm.weight": "model-00084-of-000163.safetensors", + "model.layers.33.self_attn.q_a_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.self_attn.q_a_layernorm.weight": "model-00084-of-000163.safetensors", + "model.layers.33.self_attn.q_b_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.self_attn.kv_a_proj_with_mqa.weight": "model-00084-of-000163.safetensors", + "model.layers.33.self_attn.kv_a_layernorm.weight": "model-00084-of-000163.safetensors", + "model.layers.33.self_attn.kv_b_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.self_attn.o_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.gate.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.gate.e_score_correction_bias": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.shared_experts.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.shared_experts.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.shared_experts.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.0.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.0.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.0.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.1.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.1.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.1.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.2.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.2.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.2.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.3.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.3.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.3.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.4.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.4.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.4.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.5.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.5.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.5.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.6.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.6.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.6.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.7.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.7.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.7.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.8.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.8.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.8.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.9.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.9.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.9.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.10.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.10.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.10.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.11.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.11.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.11.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.12.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.12.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.12.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.13.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.13.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.13.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.14.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.14.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.14.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.15.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.15.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.15.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.16.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.16.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.16.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.17.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.17.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.17.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.18.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.18.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.18.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.19.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.19.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.19.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.20.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.20.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.20.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.21.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.21.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.21.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.22.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.22.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.22.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.23.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.23.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.23.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.24.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.24.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.24.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.25.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.25.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.25.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.26.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.26.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.26.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.27.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.27.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.27.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.28.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.28.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.28.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.29.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.29.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.29.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.30.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.30.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.30.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.31.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.31.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.31.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.32.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.32.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.32.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.33.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.33.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.33.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.34.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.34.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.34.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.35.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.35.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.35.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.36.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.36.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.36.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.37.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.37.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.37.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.38.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.38.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.38.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.39.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.39.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.39.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.40.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.40.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.40.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.41.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.41.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.41.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.42.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.42.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.42.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.43.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.43.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.43.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.44.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.44.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.44.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.45.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.45.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.45.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.46.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.46.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.46.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.47.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.47.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.47.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.48.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.48.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.48.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.49.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.49.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.49.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.50.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.50.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.50.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.51.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.51.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.51.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.52.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.52.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.52.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.53.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.53.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.53.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.54.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.54.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.54.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.55.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.55.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.55.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.56.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.56.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.56.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.57.gate_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.57.up_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.57.down_proj.weight": "model-00084-of-000163.safetensors", + "model.layers.33.mlp.experts.58.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.58.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.58.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.59.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.59.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.59.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.60.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.60.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.60.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.61.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.61.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.61.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.62.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.62.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.62.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.63.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.63.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.63.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.64.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.64.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.64.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.65.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.65.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.65.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.66.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.66.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.66.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.67.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.67.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.67.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.68.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.68.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.68.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.69.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.69.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.69.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.70.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.70.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.70.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.71.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.71.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.71.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.72.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.72.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.72.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.73.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.73.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.73.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.74.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.74.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.74.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.75.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.75.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.75.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.76.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.76.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.76.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.77.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.77.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.77.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.78.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.78.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.78.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.79.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.79.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.79.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.80.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.80.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.80.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.81.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.81.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.81.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.82.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.82.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.82.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.83.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.83.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.83.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.84.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.84.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.84.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.85.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.85.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.85.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.86.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.86.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.86.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.87.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.87.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.87.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.88.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.88.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.88.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.89.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.89.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.89.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.90.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.90.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.90.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.91.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.91.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.91.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.92.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.92.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.92.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.93.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.93.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.93.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.94.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.94.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.94.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.95.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.95.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.95.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.96.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.96.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.96.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.97.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.97.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.97.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.98.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.98.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.98.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.99.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.99.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.99.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.100.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.100.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.100.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.101.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.101.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.101.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.102.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.102.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.102.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.103.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.103.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.103.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.104.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.104.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.104.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.105.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.105.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.105.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.106.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.106.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.106.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.107.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.107.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.107.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.108.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.108.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.108.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.109.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.109.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.109.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.110.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.110.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.110.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.111.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.111.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.111.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.112.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.112.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.112.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.113.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.113.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.113.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.114.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.114.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.114.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.115.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.115.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.115.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.116.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.116.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.116.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.117.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.117.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.117.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.118.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.118.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.118.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.119.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.119.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.119.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.120.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.120.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.120.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.121.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.121.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.121.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.122.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.122.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.122.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.123.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.123.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.123.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.124.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.124.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.124.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.125.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.125.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.125.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.126.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.126.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.126.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.127.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.127.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.127.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.128.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.128.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.128.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.129.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.129.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.129.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.130.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.130.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.130.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.131.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.131.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.131.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.132.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.132.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.132.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.133.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.133.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.133.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.134.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.134.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.134.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.135.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.135.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.135.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.136.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.136.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.136.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.137.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.137.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.137.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.138.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.138.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.138.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.139.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.139.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.139.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.140.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.140.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.140.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.141.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.141.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.141.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.142.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.142.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.142.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.143.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.143.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.143.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.144.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.144.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.144.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.145.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.145.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.145.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.146.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.146.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.146.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.147.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.147.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.147.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.148.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.148.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.148.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.149.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.149.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.149.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.150.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.150.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.150.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.151.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.151.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.151.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.152.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.152.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.152.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.153.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.153.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.153.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.154.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.154.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.154.down_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.155.gate_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.155.up_proj.weight": "model-00085-of-000163.safetensors", + "model.layers.33.mlp.experts.155.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.156.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.156.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.156.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.157.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.157.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.157.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.158.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.158.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.158.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.159.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.159.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.159.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.160.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.160.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.160.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.161.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.161.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.161.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.162.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.162.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.162.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.163.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.163.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.163.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.164.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.164.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.164.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.165.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.165.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.165.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.166.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.166.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.166.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.167.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.167.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.167.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.168.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.168.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.168.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.169.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.169.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.169.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.170.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.170.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.170.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.171.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.171.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.171.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.172.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.172.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.172.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.173.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.173.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.173.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.174.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.174.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.174.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.175.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.175.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.175.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.176.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.176.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.176.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.177.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.177.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.177.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.178.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.178.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.178.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.179.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.179.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.179.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.180.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.180.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.180.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.181.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.181.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.181.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.182.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.182.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.182.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.183.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.183.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.183.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.184.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.184.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.184.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.185.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.185.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.185.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.186.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.186.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.186.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.187.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.187.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.187.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.188.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.188.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.188.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.189.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.189.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.189.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.190.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.190.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.190.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.191.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.191.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.191.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.192.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.192.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.192.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.193.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.193.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.193.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.194.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.194.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.194.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.195.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.195.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.195.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.196.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.196.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.196.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.197.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.197.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.197.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.198.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.198.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.198.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.199.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.199.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.199.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.200.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.200.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.200.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.201.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.201.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.201.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.202.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.202.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.202.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.203.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.203.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.203.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.204.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.204.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.204.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.205.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.205.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.205.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.206.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.206.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.206.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.207.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.207.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.207.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.208.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.208.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.208.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.209.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.209.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.209.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.210.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.210.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.210.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.211.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.211.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.211.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.212.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.212.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.212.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.213.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.213.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.213.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.214.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.214.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.214.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.215.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.215.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.215.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.216.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.216.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.216.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.217.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.217.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.217.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.218.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.218.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.218.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.219.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.219.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.219.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.220.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.220.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.220.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.221.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.221.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.221.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.222.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.222.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.222.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.223.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.223.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.223.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.224.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.224.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.224.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.225.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.225.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.225.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.226.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.226.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.226.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.227.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.227.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.227.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.228.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.228.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.228.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.229.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.229.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.229.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.230.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.230.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.230.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.231.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.231.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.231.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.232.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.232.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.232.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.233.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.233.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.233.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.234.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.234.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.234.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.235.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.235.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.235.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.236.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.236.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.236.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.237.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.237.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.237.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.238.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.238.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.238.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.239.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.239.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.239.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.240.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.240.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.240.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.241.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.241.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.241.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.242.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.242.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.242.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.243.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.243.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.243.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.244.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.244.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.244.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.245.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.245.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.245.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.246.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.246.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.246.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.247.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.247.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.247.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.248.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.248.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.248.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.249.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.249.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.249.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.250.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.250.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.250.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.251.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.251.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.251.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.252.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.252.up_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.252.down_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.253.gate_proj.weight": "model-00086-of-000163.safetensors", + "model.layers.33.mlp.experts.253.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.33.mlp.experts.253.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.33.mlp.experts.254.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.33.mlp.experts.254.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.33.mlp.experts.254.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.33.mlp.experts.255.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.33.mlp.experts.255.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.33.mlp.experts.255.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.33.input_layernorm.weight": "model-00087-of-000163.safetensors", + "model.layers.33.post_attention_layernorm.weight": "model-00087-of-000163.safetensors", + "model.layers.34.self_attn.q_a_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.self_attn.q_a_layernorm.weight": "model-00087-of-000163.safetensors", + "model.layers.34.self_attn.q_b_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.self_attn.kv_a_proj_with_mqa.weight": "model-00087-of-000163.safetensors", + "model.layers.34.self_attn.kv_a_layernorm.weight": "model-00087-of-000163.safetensors", + "model.layers.34.self_attn.kv_b_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.self_attn.o_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.gate.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.gate.e_score_correction_bias": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.shared_experts.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.shared_experts.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.shared_experts.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.0.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.0.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.0.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.1.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.1.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.1.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.2.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.2.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.2.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.3.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.3.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.3.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.4.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.4.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.4.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.5.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.5.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.5.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.6.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.6.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.6.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.7.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.7.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.7.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.8.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.8.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.8.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.9.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.9.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.9.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.10.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.10.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.10.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.11.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.11.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.11.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.12.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.12.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.12.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.13.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.13.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.13.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.14.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.14.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.14.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.15.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.15.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.15.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.16.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.16.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.16.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.17.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.17.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.17.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.18.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.18.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.18.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.19.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.19.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.19.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.20.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.20.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.20.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.21.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.21.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.21.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.22.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.22.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.22.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.23.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.23.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.23.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.24.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.24.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.24.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.25.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.25.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.25.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.26.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.26.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.26.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.27.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.27.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.27.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.28.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.28.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.28.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.29.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.29.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.29.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.30.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.30.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.30.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.31.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.31.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.31.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.32.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.32.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.32.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.33.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.33.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.33.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.34.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.34.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.34.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.35.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.35.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.35.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.36.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.36.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.36.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.37.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.37.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.37.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.38.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.38.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.38.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.39.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.39.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.39.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.40.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.40.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.40.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.41.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.41.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.41.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.42.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.42.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.42.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.43.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.43.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.43.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.44.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.44.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.44.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.45.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.45.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.45.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.46.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.46.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.46.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.47.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.47.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.47.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.48.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.48.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.48.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.49.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.49.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.49.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.50.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.50.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.50.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.51.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.51.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.51.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.52.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.52.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.52.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.53.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.53.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.53.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.54.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.54.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.54.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.55.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.55.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.55.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.56.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.56.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.56.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.57.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.57.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.57.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.58.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.58.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.58.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.59.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.59.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.59.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.60.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.60.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.60.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.61.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.61.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.61.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.62.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.62.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.62.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.63.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.63.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.63.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.64.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.64.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.64.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.65.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.65.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.65.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.66.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.66.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.66.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.67.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.67.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.67.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.68.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.68.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.68.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.69.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.69.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.69.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.70.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.70.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.70.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.71.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.71.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.71.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.72.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.72.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.72.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.73.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.73.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.73.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.74.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.74.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.74.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.75.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.75.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.75.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.76.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.76.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.76.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.77.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.77.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.77.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.78.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.78.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.78.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.79.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.79.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.79.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.80.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.80.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.80.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.81.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.81.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.81.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.82.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.82.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.82.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.83.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.83.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.83.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.84.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.84.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.84.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.85.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.85.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.85.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.86.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.86.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.86.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.87.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.87.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.87.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.88.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.88.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.88.down_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.89.gate_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.89.up_proj.weight": "model-00087-of-000163.safetensors", + "model.layers.34.mlp.experts.89.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.90.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.90.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.90.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.91.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.91.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.91.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.92.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.92.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.92.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.93.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.93.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.93.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.94.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.94.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.94.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.95.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.95.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.95.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.96.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.96.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.96.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.97.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.97.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.97.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.98.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.98.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.98.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.99.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.99.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.99.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.100.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.100.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.100.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.101.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.101.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.101.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.102.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.102.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.102.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.103.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.103.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.103.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.104.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.104.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.104.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.105.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.105.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.105.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.106.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.106.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.106.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.107.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.107.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.107.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.108.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.108.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.108.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.109.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.109.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.109.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.110.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.110.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.110.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.111.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.111.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.111.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.112.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.112.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.112.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.113.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.113.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.113.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.114.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.114.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.114.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.115.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.115.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.115.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.116.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.116.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.116.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.117.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.117.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.117.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.118.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.118.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.118.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.119.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.119.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.119.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.120.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.120.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.120.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.121.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.121.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.121.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.122.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.122.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.122.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.123.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.123.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.123.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.124.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.124.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.124.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.125.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.125.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.125.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.126.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.126.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.126.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.127.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.127.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.127.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.128.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.128.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.128.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.129.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.129.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.129.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.130.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.130.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.130.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.131.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.131.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.131.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.132.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.132.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.132.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.133.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.133.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.133.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.134.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.134.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.134.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.135.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.135.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.135.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.136.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.136.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.136.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.137.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.137.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.137.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.138.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.138.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.138.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.139.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.139.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.139.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.140.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.140.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.140.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.141.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.141.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.141.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.142.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.142.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.142.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.143.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.143.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.143.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.144.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.144.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.144.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.145.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.145.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.145.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.146.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.146.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.146.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.147.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.147.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.147.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.148.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.148.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.148.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.149.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.149.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.149.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.150.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.150.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.150.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.151.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.151.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.151.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.152.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.152.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.152.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.153.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.153.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.153.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.154.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.154.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.154.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.155.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.155.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.155.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.156.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.156.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.156.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.157.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.157.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.157.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.158.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.158.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.158.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.159.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.159.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.159.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.160.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.160.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.160.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.161.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.161.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.161.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.162.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.162.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.162.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.163.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.163.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.163.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.164.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.164.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.164.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.165.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.165.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.165.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.166.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.166.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.166.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.167.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.167.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.167.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.168.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.168.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.168.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.169.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.169.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.169.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.170.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.170.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.170.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.171.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.171.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.171.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.172.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.172.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.172.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.173.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.173.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.173.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.174.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.174.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.174.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.175.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.175.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.175.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.176.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.176.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.176.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.177.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.177.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.177.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.178.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.178.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.178.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.179.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.179.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.179.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.180.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.180.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.180.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.181.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.181.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.181.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.182.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.182.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.182.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.183.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.183.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.183.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.184.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.184.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.184.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.185.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.185.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.185.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.186.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.186.up_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.186.down_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.187.gate_proj.weight": "model-00088-of-000163.safetensors", + "model.layers.34.mlp.experts.187.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.187.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.188.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.188.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.188.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.189.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.189.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.189.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.190.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.190.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.190.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.191.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.191.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.191.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.192.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.192.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.192.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.193.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.193.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.193.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.194.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.194.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.194.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.195.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.195.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.195.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.196.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.196.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.196.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.197.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.197.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.197.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.198.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.198.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.198.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.199.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.199.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.199.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.200.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.200.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.200.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.201.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.201.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.201.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.202.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.202.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.202.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.203.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.203.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.203.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.204.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.204.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.204.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.205.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.205.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.205.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.206.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.206.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.206.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.207.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.207.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.207.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.208.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.208.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.208.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.209.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.209.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.209.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.210.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.210.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.210.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.211.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.211.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.211.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.212.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.212.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.212.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.213.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.213.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.213.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.214.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.214.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.214.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.215.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.215.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.215.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.216.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.216.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.216.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.217.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.217.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.217.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.218.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.218.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.218.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.219.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.219.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.219.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.220.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.220.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.220.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.221.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.221.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.221.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.222.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.222.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.222.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.223.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.223.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.223.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.224.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.224.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.224.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.225.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.225.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.225.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.226.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.226.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.226.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.227.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.227.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.227.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.228.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.228.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.228.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.229.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.229.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.229.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.230.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.230.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.230.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.231.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.231.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.231.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.232.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.232.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.232.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.233.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.233.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.233.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.234.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.234.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.234.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.235.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.235.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.235.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.236.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.236.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.236.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.237.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.237.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.237.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.238.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.238.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.238.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.239.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.239.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.239.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.240.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.240.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.240.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.241.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.241.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.241.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.242.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.242.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.242.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.243.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.243.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.243.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.244.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.244.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.244.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.245.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.245.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.245.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.246.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.246.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.246.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.247.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.247.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.247.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.248.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.248.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.248.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.249.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.249.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.249.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.250.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.250.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.250.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.251.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.251.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.251.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.252.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.252.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.252.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.253.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.253.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.253.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.254.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.254.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.254.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.255.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.255.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.mlp.experts.255.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.34.input_layernorm.weight": "model-00089-of-000163.safetensors", + "model.layers.34.post_attention_layernorm.weight": "model-00089-of-000163.safetensors", + "model.layers.35.self_attn.q_a_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.self_attn.q_a_layernorm.weight": "model-00089-of-000163.safetensors", + "model.layers.35.self_attn.q_b_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.self_attn.kv_a_proj_with_mqa.weight": "model-00089-of-000163.safetensors", + "model.layers.35.self_attn.kv_a_layernorm.weight": "model-00089-of-000163.safetensors", + "model.layers.35.self_attn.kv_b_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.self_attn.o_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.gate.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.gate.e_score_correction_bias": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.shared_experts.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.shared_experts.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.shared_experts.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.0.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.0.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.0.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.1.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.1.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.1.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.2.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.2.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.2.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.3.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.3.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.3.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.4.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.4.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.4.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.5.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.5.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.5.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.6.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.6.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.6.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.7.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.7.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.7.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.8.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.8.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.8.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.9.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.9.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.9.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.10.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.10.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.10.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.11.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.11.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.11.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.12.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.12.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.12.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.13.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.13.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.13.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.14.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.14.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.14.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.15.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.15.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.15.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.16.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.16.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.16.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.17.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.17.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.17.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.18.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.18.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.18.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.19.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.19.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.19.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.20.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.20.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.20.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.21.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.21.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.21.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.22.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.22.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.22.down_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.23.gate_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.23.up_proj.weight": "model-00089-of-000163.safetensors", + "model.layers.35.mlp.experts.23.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.24.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.24.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.24.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.25.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.25.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.25.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.26.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.26.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.26.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.27.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.27.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.27.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.28.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.28.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.28.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.29.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.29.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.29.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.30.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.30.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.30.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.31.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.31.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.31.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.32.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.32.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.32.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.33.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.33.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.33.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.34.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.34.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.34.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.35.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.35.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.35.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.36.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.36.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.36.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.37.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.37.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.37.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.38.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.38.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.38.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.39.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.39.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.39.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.40.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.40.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.40.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.41.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.41.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.41.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.42.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.42.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.42.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.43.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.43.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.43.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.44.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.44.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.44.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.45.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.45.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.45.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.46.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.46.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.46.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.47.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.47.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.47.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.48.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.48.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.48.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.49.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.49.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.49.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.50.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.50.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.50.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.51.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.51.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.51.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.52.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.52.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.52.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.53.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.53.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.53.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.54.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.54.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.54.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.55.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.55.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.55.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.56.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.56.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.56.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.57.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.57.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.57.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.58.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.58.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.58.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.59.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.59.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.59.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.60.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.60.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.60.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.61.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.61.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.61.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.62.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.62.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.62.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.63.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.63.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.63.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.64.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.64.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.64.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.65.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.65.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.65.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.66.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.66.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.66.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.67.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.67.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.67.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.68.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.68.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.68.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.69.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.69.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.69.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.70.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.70.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.70.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.71.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.71.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.71.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.72.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.72.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.72.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.73.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.73.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.73.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.74.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.74.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.74.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.75.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.75.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.75.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.76.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.76.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.76.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.77.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.77.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.77.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.78.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.78.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.78.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.79.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.79.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.79.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.80.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.80.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.80.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.81.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.81.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.81.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.82.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.82.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.82.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.83.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.83.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.83.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.84.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.84.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.84.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.85.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.85.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.85.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.86.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.86.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.86.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.87.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.87.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.87.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.88.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.88.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.88.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.89.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.89.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.89.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.90.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.90.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.90.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.91.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.91.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.91.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.92.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.92.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.92.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.93.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.93.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.93.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.94.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.94.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.94.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.95.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.95.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.95.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.96.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.96.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.96.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.97.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.97.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.97.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.98.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.98.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.98.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.99.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.99.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.99.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.100.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.100.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.100.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.101.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.101.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.101.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.102.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.102.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.102.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.103.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.103.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.103.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.104.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.104.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.104.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.105.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.105.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.105.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.106.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.106.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.106.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.107.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.107.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.107.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.108.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.108.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.108.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.109.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.109.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.109.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.110.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.110.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.110.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.111.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.111.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.111.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.112.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.112.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.112.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.113.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.113.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.113.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.114.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.114.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.114.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.115.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.115.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.115.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.116.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.116.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.116.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.117.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.117.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.117.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.118.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.118.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.118.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.119.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.119.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.119.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.120.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.120.up_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.120.down_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.121.gate_proj.weight": "model-00090-of-000163.safetensors", + "model.layers.35.mlp.experts.121.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.121.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.122.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.122.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.122.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.123.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.123.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.123.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.124.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.124.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.124.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.125.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.125.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.125.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.126.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.126.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.126.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.127.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.127.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.127.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.128.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.128.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.128.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.129.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.129.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.129.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.130.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.130.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.130.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.131.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.131.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.131.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.132.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.132.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.132.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.133.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.133.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.133.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.134.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.134.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.134.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.135.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.135.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.135.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.136.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.136.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.136.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.137.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.137.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.137.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.138.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.138.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.138.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.139.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.139.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.139.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.140.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.140.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.140.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.141.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.141.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.141.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.142.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.142.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.142.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.143.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.143.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.143.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.144.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.144.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.144.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.145.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.145.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.145.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.146.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.146.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.146.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.147.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.147.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.147.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.148.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.148.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.148.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.149.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.149.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.149.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.150.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.150.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.150.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.151.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.151.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.151.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.152.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.152.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.152.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.153.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.153.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.153.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.154.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.154.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.154.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.155.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.155.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.155.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.156.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.156.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.156.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.157.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.157.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.157.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.158.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.158.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.158.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.159.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.159.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.159.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.160.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.160.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.160.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.161.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.161.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.161.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.162.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.162.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.162.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.163.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.163.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.163.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.164.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.164.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.164.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.165.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.165.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.165.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.166.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.166.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.166.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.167.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.167.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.167.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.168.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.168.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.168.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.169.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.169.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.169.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.170.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.170.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.170.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.171.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.171.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.171.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.172.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.172.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.172.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.173.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.173.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.173.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.174.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.174.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.174.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.175.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.175.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.175.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.176.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.176.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.176.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.177.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.177.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.177.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.178.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.178.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.178.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.179.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.179.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.179.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.180.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.180.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.180.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.181.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.181.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.181.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.182.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.182.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.182.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.183.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.183.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.183.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.184.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.184.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.184.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.185.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.185.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.185.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.186.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.186.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.186.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.187.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.187.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.187.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.188.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.188.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.188.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.189.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.189.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.189.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.190.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.190.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.190.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.191.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.191.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.191.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.192.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.192.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.192.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.193.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.193.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.193.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.194.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.194.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.194.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.195.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.195.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.195.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.196.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.196.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.196.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.197.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.197.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.197.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.198.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.198.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.198.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.199.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.199.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.199.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.200.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.200.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.200.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.201.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.201.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.201.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.202.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.202.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.202.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.203.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.203.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.203.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.204.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.204.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.204.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.205.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.205.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.205.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.206.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.206.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.206.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.207.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.207.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.207.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.208.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.208.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.208.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.209.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.209.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.209.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.210.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.210.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.210.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.211.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.211.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.211.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.212.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.212.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.212.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.213.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.213.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.213.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.214.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.214.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.214.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.215.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.215.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.215.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.216.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.216.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.216.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.217.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.217.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.217.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.218.gate_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.218.up_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.218.down_proj.weight": "model-00091-of-000163.safetensors", + "model.layers.35.mlp.experts.219.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.219.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.219.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.220.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.220.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.220.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.221.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.221.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.221.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.222.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.222.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.222.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.223.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.223.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.223.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.224.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.224.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.224.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.225.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.225.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.225.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.226.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.226.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.226.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.227.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.227.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.227.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.228.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.228.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.228.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.229.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.229.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.229.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.230.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.230.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.230.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.231.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.231.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.231.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.232.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.232.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.232.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.233.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.233.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.233.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.234.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.234.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.234.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.235.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.235.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.235.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.236.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.236.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.236.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.237.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.237.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.237.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.238.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.238.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.238.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.239.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.239.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.239.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.240.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.240.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.240.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.241.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.241.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.241.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.242.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.242.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.242.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.243.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.243.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.243.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.244.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.244.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.244.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.245.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.245.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.245.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.246.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.246.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.246.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.247.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.247.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.247.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.248.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.248.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.248.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.249.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.249.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.249.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.250.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.250.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.250.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.251.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.251.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.251.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.252.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.252.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.252.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.253.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.253.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.253.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.254.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.254.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.254.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.255.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.255.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.mlp.experts.255.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.35.input_layernorm.weight": "model-00092-of-000163.safetensors", + "model.layers.35.post_attention_layernorm.weight": "model-00092-of-000163.safetensors", + "model.layers.36.self_attn.q_a_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.self_attn.q_a_layernorm.weight": "model-00092-of-000163.safetensors", + "model.layers.36.self_attn.q_b_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.self_attn.kv_a_proj_with_mqa.weight": "model-00092-of-000163.safetensors", + "model.layers.36.self_attn.kv_a_layernorm.weight": "model-00092-of-000163.safetensors", + "model.layers.36.self_attn.kv_b_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.self_attn.o_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.gate.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.gate.e_score_correction_bias": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.shared_experts.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.shared_experts.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.shared_experts.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.0.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.0.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.0.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.1.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.1.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.1.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.2.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.2.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.2.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.3.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.3.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.3.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.4.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.4.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.4.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.5.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.5.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.5.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.6.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.6.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.6.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.7.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.7.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.7.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.8.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.8.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.8.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.9.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.9.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.9.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.10.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.10.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.10.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.11.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.11.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.11.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.12.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.12.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.12.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.13.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.13.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.13.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.14.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.14.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.14.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.15.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.15.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.15.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.16.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.16.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.16.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.17.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.17.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.17.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.18.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.18.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.18.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.19.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.19.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.19.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.20.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.20.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.20.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.21.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.21.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.21.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.22.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.22.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.22.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.23.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.23.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.23.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.24.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.24.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.24.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.25.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.25.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.25.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.26.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.26.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.26.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.27.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.27.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.27.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.28.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.28.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.28.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.29.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.29.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.29.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.30.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.30.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.30.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.31.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.31.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.31.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.32.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.32.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.32.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.33.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.33.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.33.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.34.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.34.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.34.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.35.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.35.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.35.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.36.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.36.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.36.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.37.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.37.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.37.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.38.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.38.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.38.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.39.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.39.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.39.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.40.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.40.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.40.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.41.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.41.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.41.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.42.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.42.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.42.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.43.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.43.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.43.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.44.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.44.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.44.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.45.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.45.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.45.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.46.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.46.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.46.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.47.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.47.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.47.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.48.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.48.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.48.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.49.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.49.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.49.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.50.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.50.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.50.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.51.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.51.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.51.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.52.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.52.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.52.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.53.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.53.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.53.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.54.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.54.up_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.54.down_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.55.gate_proj.weight": "model-00092-of-000163.safetensors", + "model.layers.36.mlp.experts.55.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.55.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.56.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.56.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.56.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.57.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.57.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.57.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.58.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.58.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.58.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.59.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.59.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.59.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.60.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.60.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.60.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.61.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.61.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.61.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.62.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.62.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.62.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.63.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.63.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.63.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.64.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.64.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.64.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.65.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.65.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.65.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.66.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.66.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.66.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.67.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.67.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.67.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.68.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.68.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.68.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.69.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.69.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.69.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.70.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.70.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.70.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.71.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.71.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.71.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.72.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.72.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.72.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.73.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.73.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.73.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.74.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.74.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.74.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.75.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.75.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.75.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.76.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.76.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.76.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.77.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.77.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.77.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.78.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.78.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.78.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.79.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.79.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.79.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.80.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.80.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.80.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.81.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.81.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.81.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.82.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.82.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.82.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.83.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.83.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.83.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.84.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.84.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.84.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.85.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.85.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.85.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.86.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.86.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.86.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.87.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.87.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.87.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.88.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.88.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.88.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.89.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.89.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.89.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.90.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.90.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.90.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.91.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.91.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.91.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.92.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.92.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.92.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.93.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.93.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.93.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.94.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.94.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.94.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.95.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.95.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.95.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.96.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.96.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.96.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.97.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.97.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.97.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.98.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.98.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.98.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.99.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.99.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.99.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.100.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.100.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.100.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.101.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.101.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.101.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.102.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.102.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.102.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.103.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.103.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.103.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.104.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.104.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.104.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.105.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.105.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.105.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.106.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.106.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.106.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.107.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.107.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.107.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.108.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.108.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.108.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.109.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.109.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.109.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.110.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.110.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.110.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.111.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.111.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.111.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.112.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.112.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.112.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.113.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.113.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.113.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.114.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.114.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.114.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.115.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.115.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.115.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.116.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.116.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.116.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.117.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.117.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.117.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.118.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.118.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.118.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.119.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.119.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.119.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.120.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.120.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.120.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.121.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.121.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.121.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.122.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.122.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.122.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.123.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.123.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.123.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.124.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.124.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.124.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.125.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.125.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.125.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.126.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.126.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.126.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.127.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.127.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.127.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.128.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.128.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.128.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.129.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.129.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.129.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.130.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.130.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.130.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.131.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.131.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.131.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.132.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.132.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.132.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.133.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.133.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.133.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.134.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.134.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.134.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.135.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.135.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.135.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.136.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.136.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.136.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.137.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.137.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.137.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.138.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.138.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.138.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.139.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.139.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.139.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.140.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.140.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.140.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.141.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.141.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.141.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.142.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.142.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.142.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.143.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.143.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.143.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.144.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.144.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.144.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.145.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.145.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.145.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.146.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.146.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.146.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.147.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.147.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.147.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.148.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.148.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.148.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.149.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.149.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.149.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.150.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.150.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.150.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.151.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.151.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.151.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.152.gate_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.152.up_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.152.down_proj.weight": "model-00093-of-000163.safetensors", + "model.layers.36.mlp.experts.153.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.153.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.153.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.154.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.154.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.154.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.155.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.155.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.155.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.156.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.156.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.156.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.157.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.157.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.157.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.158.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.158.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.158.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.159.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.159.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.159.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.160.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.160.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.160.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.161.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.161.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.161.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.162.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.162.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.162.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.163.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.163.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.163.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.164.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.164.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.164.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.165.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.165.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.165.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.166.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.166.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.166.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.167.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.167.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.167.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.168.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.168.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.168.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.169.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.169.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.169.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.170.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.170.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.170.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.171.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.171.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.171.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.172.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.172.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.172.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.173.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.173.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.173.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.174.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.174.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.174.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.175.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.175.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.175.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.176.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.176.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.176.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.177.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.177.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.177.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.178.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.178.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.178.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.179.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.179.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.179.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.180.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.180.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.180.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.181.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.181.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.181.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.182.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.182.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.182.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.183.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.183.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.183.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.184.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.184.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.184.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.185.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.185.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.185.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.186.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.186.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.186.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.187.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.187.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.187.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.188.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.188.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.188.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.189.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.189.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.189.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.190.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.190.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.190.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.191.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.191.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.191.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.192.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.192.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.192.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.193.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.193.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.193.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.194.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.194.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.194.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.195.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.195.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.195.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.196.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.196.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.196.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.197.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.197.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.197.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.198.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.198.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.198.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.199.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.199.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.199.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.200.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.200.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.200.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.201.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.201.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.201.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.202.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.202.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.202.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.203.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.203.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.203.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.204.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.204.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.204.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.205.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.205.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.205.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.206.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.206.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.206.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.207.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.207.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.207.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.208.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.208.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.208.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.209.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.209.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.209.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.210.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.210.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.210.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.211.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.211.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.211.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.212.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.212.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.212.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.213.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.213.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.213.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.214.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.214.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.214.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.215.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.215.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.215.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.216.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.216.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.216.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.217.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.217.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.217.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.218.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.218.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.218.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.219.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.219.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.219.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.220.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.220.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.220.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.221.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.221.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.221.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.222.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.222.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.222.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.223.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.223.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.223.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.224.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.224.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.224.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.225.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.225.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.225.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.226.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.226.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.226.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.227.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.227.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.227.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.228.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.228.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.228.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.229.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.229.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.229.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.230.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.230.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.230.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.231.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.231.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.231.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.232.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.232.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.232.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.233.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.233.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.233.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.234.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.234.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.234.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.235.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.235.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.235.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.236.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.236.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.236.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.237.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.237.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.237.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.238.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.238.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.238.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.239.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.239.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.239.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.240.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.240.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.240.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.241.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.241.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.241.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.242.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.242.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.242.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.243.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.243.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.243.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.244.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.244.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.244.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.245.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.245.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.245.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.246.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.246.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.246.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.247.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.247.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.247.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.248.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.248.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.248.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.249.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.249.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.249.down_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.250.gate_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.250.up_proj.weight": "model-00094-of-000163.safetensors", + "model.layers.36.mlp.experts.250.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.36.mlp.experts.251.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.36.mlp.experts.251.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.36.mlp.experts.251.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.36.mlp.experts.252.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.36.mlp.experts.252.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.36.mlp.experts.252.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.36.mlp.experts.253.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.36.mlp.experts.253.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.36.mlp.experts.253.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.36.mlp.experts.254.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.36.mlp.experts.254.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.36.mlp.experts.254.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.36.mlp.experts.255.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.36.mlp.experts.255.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.36.mlp.experts.255.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.36.input_layernorm.weight": "model-00095-of-000163.safetensors", + "model.layers.36.post_attention_layernorm.weight": "model-00095-of-000163.safetensors", + "model.layers.37.self_attn.q_a_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.self_attn.q_a_layernorm.weight": "model-00095-of-000163.safetensors", + "model.layers.37.self_attn.q_b_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.self_attn.kv_a_proj_with_mqa.weight": "model-00095-of-000163.safetensors", + "model.layers.37.self_attn.kv_a_layernorm.weight": "model-00095-of-000163.safetensors", + "model.layers.37.self_attn.kv_b_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.self_attn.o_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.gate.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.gate.e_score_correction_bias": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.shared_experts.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.shared_experts.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.shared_experts.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.0.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.0.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.0.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.1.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.1.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.1.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.2.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.2.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.2.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.3.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.3.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.3.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.4.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.4.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.4.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.5.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.5.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.5.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.6.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.6.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.6.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.7.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.7.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.7.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.8.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.8.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.8.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.9.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.9.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.9.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.10.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.10.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.10.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.11.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.11.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.11.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.12.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.12.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.12.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.13.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.13.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.13.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.14.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.14.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.14.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.15.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.15.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.15.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.16.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.16.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.16.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.17.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.17.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.17.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.18.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.18.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.18.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.19.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.19.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.19.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.20.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.20.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.20.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.21.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.21.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.21.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.22.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.22.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.22.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.23.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.23.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.23.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.24.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.24.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.24.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.25.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.25.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.25.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.26.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.26.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.26.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.27.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.27.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.27.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.28.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.28.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.28.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.29.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.29.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.29.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.30.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.30.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.30.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.31.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.31.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.31.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.32.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.32.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.32.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.33.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.33.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.33.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.34.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.34.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.34.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.35.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.35.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.35.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.36.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.36.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.36.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.37.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.37.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.37.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.38.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.38.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.38.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.39.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.39.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.39.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.40.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.40.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.40.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.41.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.41.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.41.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.42.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.42.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.42.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.43.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.43.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.43.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.44.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.44.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.44.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.45.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.45.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.45.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.46.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.46.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.46.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.47.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.47.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.47.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.48.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.48.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.48.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.49.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.49.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.49.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.50.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.50.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.50.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.51.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.51.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.51.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.52.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.52.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.52.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.53.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.53.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.53.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.54.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.54.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.54.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.55.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.55.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.55.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.56.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.56.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.56.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.57.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.57.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.57.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.58.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.58.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.58.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.59.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.59.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.59.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.60.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.60.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.60.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.61.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.61.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.61.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.62.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.62.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.62.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.63.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.63.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.63.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.64.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.64.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.64.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.65.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.65.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.65.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.66.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.66.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.66.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.67.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.67.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.67.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.68.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.68.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.68.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.69.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.69.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.69.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.70.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.70.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.70.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.71.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.71.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.71.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.72.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.72.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.72.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.73.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.73.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.73.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.74.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.74.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.74.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.75.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.75.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.75.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.76.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.76.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.76.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.77.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.77.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.77.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.78.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.78.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.78.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.79.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.79.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.79.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.80.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.80.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.80.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.81.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.81.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.81.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.82.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.82.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.82.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.83.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.83.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.83.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.84.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.84.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.84.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.85.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.85.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.85.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.86.gate_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.86.up_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.86.down_proj.weight": "model-00095-of-000163.safetensors", + "model.layers.37.mlp.experts.87.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.87.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.87.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.88.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.88.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.88.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.89.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.89.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.89.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.90.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.90.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.90.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.91.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.91.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.91.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.92.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.92.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.92.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.93.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.93.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.93.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.94.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.94.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.94.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.95.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.95.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.95.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.96.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.96.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.96.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.97.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.97.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.97.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.98.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.98.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.98.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.99.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.99.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.99.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.100.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.100.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.100.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.101.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.101.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.101.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.102.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.102.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.102.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.103.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.103.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.103.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.104.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.104.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.104.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.105.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.105.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.105.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.106.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.106.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.106.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.107.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.107.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.107.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.108.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.108.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.108.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.109.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.109.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.109.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.110.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.110.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.110.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.111.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.111.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.111.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.112.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.112.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.112.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.113.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.113.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.113.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.114.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.114.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.114.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.115.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.115.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.115.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.116.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.116.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.116.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.117.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.117.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.117.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.118.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.118.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.118.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.119.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.119.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.119.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.120.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.120.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.120.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.121.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.121.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.121.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.122.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.122.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.122.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.123.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.123.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.123.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.124.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.124.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.124.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.125.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.125.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.125.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.126.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.126.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.126.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.127.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.127.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.127.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.128.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.128.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.128.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.129.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.129.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.129.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.130.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.130.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.130.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.131.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.131.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.131.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.132.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.132.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.132.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.133.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.133.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.133.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.134.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.134.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.134.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.135.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.135.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.135.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.136.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.136.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.136.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.137.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.137.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.137.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.138.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.138.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.138.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.139.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.139.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.139.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.140.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.140.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.140.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.141.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.141.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.141.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.142.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.142.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.142.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.143.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.143.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.143.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.144.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.144.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.144.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.145.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.145.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.145.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.146.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.146.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.146.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.147.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.147.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.147.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.148.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.148.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.148.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.149.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.149.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.149.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.150.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.150.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.150.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.151.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.151.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.151.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.152.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.152.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.152.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.153.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.153.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.153.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.154.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.154.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.154.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.155.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.155.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.155.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.156.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.156.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.156.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.157.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.157.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.157.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.158.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.158.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.158.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.159.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.159.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.159.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.160.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.160.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.160.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.161.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.161.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.161.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.162.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.162.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.162.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.163.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.163.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.163.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.164.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.164.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.164.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.165.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.165.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.165.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.166.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.166.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.166.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.167.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.167.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.167.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.168.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.168.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.168.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.169.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.169.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.169.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.170.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.170.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.170.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.171.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.171.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.171.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.172.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.172.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.172.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.173.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.173.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.173.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.174.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.174.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.174.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.175.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.175.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.175.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.176.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.176.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.176.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.177.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.177.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.177.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.178.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.178.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.178.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.179.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.179.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.179.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.180.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.180.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.180.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.181.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.181.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.181.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.182.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.182.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.182.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.183.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.183.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.183.down_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.184.gate_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.184.up_proj.weight": "model-00096-of-000163.safetensors", + "model.layers.37.mlp.experts.184.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.185.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.185.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.185.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.186.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.186.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.186.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.187.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.187.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.187.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.188.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.188.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.188.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.189.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.189.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.189.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.190.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.190.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.190.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.191.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.191.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.191.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.192.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.192.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.192.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.193.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.193.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.193.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.194.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.194.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.194.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.195.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.195.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.195.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.196.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.196.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.196.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.197.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.197.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.197.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.198.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.198.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.198.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.199.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.199.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.199.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.200.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.200.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.200.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.201.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.201.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.201.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.202.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.202.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.202.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.203.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.203.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.203.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.204.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.204.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.204.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.205.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.205.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.205.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.206.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.206.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.206.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.207.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.207.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.207.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.208.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.208.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.208.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.209.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.209.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.209.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.210.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.210.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.210.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.211.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.211.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.211.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.212.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.212.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.212.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.213.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.213.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.213.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.214.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.214.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.214.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.215.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.215.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.215.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.216.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.216.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.216.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.217.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.217.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.217.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.218.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.218.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.218.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.219.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.219.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.219.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.220.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.220.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.220.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.221.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.221.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.221.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.222.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.222.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.222.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.223.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.223.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.223.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.224.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.224.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.224.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.225.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.225.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.225.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.226.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.226.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.226.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.227.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.227.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.227.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.228.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.228.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.228.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.229.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.229.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.229.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.230.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.230.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.230.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.231.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.231.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.231.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.232.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.232.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.232.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.233.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.233.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.233.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.234.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.234.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.234.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.235.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.235.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.235.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.236.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.236.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.236.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.237.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.237.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.237.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.238.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.238.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.238.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.239.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.239.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.239.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.240.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.240.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.240.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.241.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.241.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.241.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.242.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.242.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.242.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.243.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.243.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.243.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.244.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.244.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.244.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.245.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.245.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.245.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.246.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.246.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.246.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.247.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.247.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.247.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.248.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.248.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.248.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.249.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.249.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.249.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.250.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.250.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.250.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.251.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.251.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.251.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.252.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.252.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.252.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.253.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.253.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.253.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.254.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.254.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.254.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.255.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.255.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.mlp.experts.255.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.37.input_layernorm.weight": "model-00097-of-000163.safetensors", + "model.layers.37.post_attention_layernorm.weight": "model-00097-of-000163.safetensors", + "model.layers.38.self_attn.q_a_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.self_attn.q_a_layernorm.weight": "model-00097-of-000163.safetensors", + "model.layers.38.self_attn.q_b_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.self_attn.kv_a_proj_with_mqa.weight": "model-00097-of-000163.safetensors", + "model.layers.38.self_attn.kv_a_layernorm.weight": "model-00097-of-000163.safetensors", + "model.layers.38.self_attn.kv_b_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.self_attn.o_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.gate.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.gate.e_score_correction_bias": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.shared_experts.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.shared_experts.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.shared_experts.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.0.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.0.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.0.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.1.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.1.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.1.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.2.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.2.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.2.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.3.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.3.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.3.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.4.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.4.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.4.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.5.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.5.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.5.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.6.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.6.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.6.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.7.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.7.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.7.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.8.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.8.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.8.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.9.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.9.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.9.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.10.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.10.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.10.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.11.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.11.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.11.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.12.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.12.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.12.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.13.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.13.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.13.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.14.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.14.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.14.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.15.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.15.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.15.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.16.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.16.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.16.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.17.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.17.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.17.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.18.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.18.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.18.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.19.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.19.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.19.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.20.gate_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.20.up_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.20.down_proj.weight": "model-00097-of-000163.safetensors", + "model.layers.38.mlp.experts.21.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.21.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.21.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.22.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.22.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.22.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.23.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.23.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.23.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.24.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.24.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.24.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.25.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.25.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.25.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.26.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.26.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.26.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.27.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.27.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.27.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.28.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.28.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.28.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.29.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.29.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.29.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.30.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.30.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.30.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.31.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.31.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.31.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.32.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.32.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.32.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.33.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.33.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.33.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.34.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.34.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.34.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.35.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.35.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.35.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.36.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.36.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.36.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.37.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.37.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.37.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.38.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.38.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.38.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.39.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.39.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.39.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.40.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.40.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.40.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.41.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.41.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.41.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.42.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.42.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.42.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.43.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.43.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.43.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.44.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.44.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.44.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.45.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.45.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.45.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.46.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.46.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.46.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.47.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.47.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.47.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.48.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.48.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.48.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.49.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.49.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.49.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.50.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.50.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.50.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.51.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.51.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.51.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.52.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.52.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.52.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.53.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.53.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.53.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.54.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.54.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.54.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.55.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.55.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.55.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.56.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.56.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.56.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.57.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.57.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.57.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.58.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.58.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.58.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.59.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.59.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.59.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.60.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.60.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.60.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.61.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.61.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.61.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.62.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.62.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.62.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.63.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.63.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.63.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.64.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.64.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.64.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.65.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.65.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.65.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.66.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.66.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.66.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.67.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.67.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.67.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.68.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.68.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.68.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.69.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.69.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.69.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.70.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.70.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.70.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.71.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.71.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.71.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.72.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.72.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.72.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.73.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.73.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.73.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.74.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.74.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.74.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.75.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.75.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.75.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.76.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.76.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.76.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.77.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.77.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.77.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.78.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.78.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.78.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.79.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.79.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.79.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.80.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.80.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.80.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.81.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.81.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.81.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.82.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.82.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.82.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.83.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.83.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.83.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.84.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.84.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.84.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.85.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.85.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.85.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.86.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.86.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.86.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.87.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.87.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.87.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.88.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.88.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.88.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.89.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.89.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.89.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.90.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.90.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.90.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.91.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.91.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.91.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.92.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.92.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.92.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.93.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.93.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.93.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.94.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.94.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.94.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.95.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.95.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.95.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.96.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.96.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.96.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.97.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.97.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.97.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.98.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.98.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.98.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.99.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.99.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.99.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.100.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.100.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.100.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.101.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.101.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.101.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.102.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.102.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.102.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.103.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.103.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.103.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.104.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.104.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.104.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.105.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.105.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.105.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.106.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.106.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.106.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.107.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.107.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.107.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.108.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.108.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.108.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.109.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.109.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.109.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.110.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.110.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.110.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.111.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.111.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.111.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.112.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.112.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.112.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.113.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.113.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.113.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.114.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.114.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.114.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.115.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.115.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.115.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.116.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.116.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.116.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.117.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.117.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.117.down_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.118.gate_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.118.up_proj.weight": "model-00098-of-000163.safetensors", + "model.layers.38.mlp.experts.118.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.119.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.119.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.119.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.120.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.120.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.120.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.121.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.121.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.121.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.122.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.122.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.122.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.123.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.123.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.123.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.124.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.124.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.124.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.125.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.125.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.125.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.126.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.126.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.126.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.127.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.127.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.127.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.128.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.128.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.128.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.129.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.129.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.129.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.130.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.130.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.130.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.131.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.131.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.131.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.132.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.132.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.132.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.133.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.133.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.133.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.134.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.134.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.134.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.135.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.135.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.135.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.136.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.136.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.136.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.137.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.137.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.137.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.138.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.138.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.138.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.139.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.139.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.139.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.140.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.140.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.140.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.141.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.141.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.141.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.142.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.142.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.142.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.143.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.143.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.143.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.144.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.144.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.144.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.145.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.145.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.145.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.146.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.146.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.146.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.147.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.147.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.147.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.148.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.148.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.148.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.149.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.149.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.149.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.150.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.150.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.150.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.151.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.151.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.151.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.152.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.152.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.152.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.153.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.153.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.153.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.154.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.154.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.154.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.155.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.155.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.155.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.156.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.156.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.156.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.157.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.157.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.157.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.158.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.158.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.158.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.159.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.159.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.159.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.160.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.160.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.160.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.161.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.161.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.161.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.162.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.162.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.162.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.163.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.163.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.163.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.164.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.164.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.164.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.165.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.165.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.165.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.166.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.166.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.166.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.167.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.167.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.167.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.168.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.168.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.168.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.169.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.169.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.169.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.170.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.170.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.170.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.171.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.171.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.171.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.172.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.172.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.172.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.173.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.173.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.173.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.174.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.174.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.174.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.175.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.175.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.175.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.176.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.176.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.176.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.177.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.177.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.177.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.178.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.178.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.178.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.179.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.179.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.179.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.180.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.180.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.180.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.181.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.181.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.181.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.182.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.182.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.182.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.183.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.183.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.183.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.184.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.184.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.184.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.185.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.185.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.185.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.186.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.186.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.186.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.187.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.187.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.187.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.188.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.188.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.188.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.189.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.189.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.189.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.190.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.190.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.190.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.191.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.191.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.191.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.192.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.192.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.192.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.193.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.193.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.193.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.194.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.194.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.194.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.195.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.195.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.195.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.196.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.196.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.196.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.197.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.197.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.197.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.198.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.198.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.198.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.199.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.199.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.199.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.200.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.200.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.200.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.201.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.201.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.201.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.202.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.202.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.202.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.203.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.203.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.203.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.204.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.204.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.204.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.205.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.205.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.205.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.206.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.206.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.206.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.207.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.207.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.207.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.208.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.208.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.208.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.209.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.209.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.209.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.210.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.210.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.210.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.211.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.211.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.211.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.212.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.212.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.212.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.213.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.213.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.213.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.214.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.214.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.214.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.215.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.215.up_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.215.down_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.216.gate_proj.weight": "model-00099-of-000163.safetensors", + "model.layers.38.mlp.experts.216.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.216.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.217.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.217.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.217.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.218.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.218.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.218.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.219.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.219.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.219.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.220.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.220.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.220.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.221.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.221.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.221.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.222.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.222.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.222.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.223.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.223.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.223.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.224.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.224.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.224.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.225.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.225.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.225.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.226.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.226.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.226.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.227.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.227.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.227.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.228.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.228.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.228.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.229.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.229.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.229.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.230.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.230.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.230.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.231.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.231.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.231.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.232.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.232.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.232.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.233.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.233.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.233.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.234.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.234.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.234.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.235.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.235.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.235.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.236.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.236.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.236.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.237.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.237.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.237.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.238.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.238.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.238.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.239.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.239.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.239.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.240.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.240.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.240.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.241.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.241.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.241.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.242.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.242.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.242.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.243.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.243.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.243.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.244.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.244.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.244.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.245.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.245.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.245.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.246.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.246.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.246.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.247.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.247.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.247.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.248.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.248.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.248.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.249.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.249.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.249.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.250.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.250.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.250.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.251.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.251.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.251.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.252.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.252.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.252.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.253.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.253.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.253.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.254.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.254.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.254.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.255.gate_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.255.up_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.mlp.experts.255.down_proj.weight": "model-00100-of-000163.safetensors", + "model.layers.38.input_layernorm.weight": "model-00100-of-000163.safetensors", + "model.layers.38.post_attention_layernorm.weight": "model-00100-of-000163.safetensors", + "model.layers.39.self_attn.q_a_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.self_attn.q_a_layernorm.weight": "model-00101-of-000163.safetensors", + "model.layers.39.self_attn.q_b_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.self_attn.kv_a_proj_with_mqa.weight": "model-00101-of-000163.safetensors", + "model.layers.39.self_attn.kv_a_layernorm.weight": "model-00101-of-000163.safetensors", + "model.layers.39.self_attn.kv_b_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.self_attn.o_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.gate.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.gate.e_score_correction_bias": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.shared_experts.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.shared_experts.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.shared_experts.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.0.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.0.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.0.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.1.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.1.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.1.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.2.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.2.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.2.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.3.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.3.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.3.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.4.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.4.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.4.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.5.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.5.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.5.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.6.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.6.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.6.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.7.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.7.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.7.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.8.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.8.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.8.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.9.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.9.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.9.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.10.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.10.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.10.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.11.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.11.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.11.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.12.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.12.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.12.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.13.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.13.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.13.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.14.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.14.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.14.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.15.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.15.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.15.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.16.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.16.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.16.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.17.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.17.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.17.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.18.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.18.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.18.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.19.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.19.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.19.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.20.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.20.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.20.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.21.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.21.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.21.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.22.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.22.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.22.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.23.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.23.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.23.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.24.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.24.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.24.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.25.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.25.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.25.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.26.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.26.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.26.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.27.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.27.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.27.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.28.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.28.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.28.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.29.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.29.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.29.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.30.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.30.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.30.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.31.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.31.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.31.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.32.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.32.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.32.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.33.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.33.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.33.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.34.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.34.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.34.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.35.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.35.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.35.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.36.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.36.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.36.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.37.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.37.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.37.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.38.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.38.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.38.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.39.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.39.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.39.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.40.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.40.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.40.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.41.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.41.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.41.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.42.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.42.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.42.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.43.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.43.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.43.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.44.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.44.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.44.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.45.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.45.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.45.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.46.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.46.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.46.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.47.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.47.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.47.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.48.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.48.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.48.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.49.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.49.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.49.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.50.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.50.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.50.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.51.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.51.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.51.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.52.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.52.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.52.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.53.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.53.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.53.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.54.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.54.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.54.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.55.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.55.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.55.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.56.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.56.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.56.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.57.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.57.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.57.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.58.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.58.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.58.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.59.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.59.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.59.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.60.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.60.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.60.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.61.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.61.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.61.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.62.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.62.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.62.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.63.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.63.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.63.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.64.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.64.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.64.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.65.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.65.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.65.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.66.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.66.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.66.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.67.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.67.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.67.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.68.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.68.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.68.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.69.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.69.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.69.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.70.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.70.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.70.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.71.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.71.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.71.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.72.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.72.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.72.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.73.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.73.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.73.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.74.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.74.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.74.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.75.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.75.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.75.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.76.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.76.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.76.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.77.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.77.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.77.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.78.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.78.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.78.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.79.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.79.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.79.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.80.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.80.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.80.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.81.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.81.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.81.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.82.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.82.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.82.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.83.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.83.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.83.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.84.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.84.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.84.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.85.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.85.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.85.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.86.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.86.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.86.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.87.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.87.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.87.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.88.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.88.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.88.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.89.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.89.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.89.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.90.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.90.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.90.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.91.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.91.up_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.91.down_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.92.gate_proj.weight": "model-00101-of-000163.safetensors", + "model.layers.39.mlp.experts.92.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.92.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.93.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.93.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.93.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.94.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.94.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.94.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.95.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.95.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.95.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.96.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.96.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.96.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.97.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.97.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.97.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.98.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.98.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.98.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.99.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.99.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.99.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.100.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.100.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.100.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.101.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.101.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.101.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.102.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.102.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.102.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.103.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.103.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.103.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.104.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.104.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.104.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.105.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.105.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.105.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.106.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.106.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.106.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.107.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.107.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.107.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.108.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.108.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.108.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.109.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.109.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.109.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.110.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.110.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.110.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.111.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.111.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.111.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.112.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.112.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.112.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.113.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.113.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.113.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.114.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.114.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.114.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.115.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.115.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.115.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.116.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.116.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.116.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.117.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.117.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.117.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.118.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.118.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.118.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.119.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.119.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.119.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.120.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.120.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.120.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.121.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.121.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.121.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.122.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.122.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.122.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.123.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.123.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.123.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.124.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.124.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.124.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.125.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.125.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.125.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.126.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.126.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.126.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.127.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.127.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.127.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.128.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.128.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.128.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.129.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.129.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.129.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.130.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.130.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.130.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.131.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.131.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.131.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.132.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.132.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.132.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.133.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.133.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.133.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.134.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.134.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.134.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.135.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.135.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.135.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.136.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.136.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.136.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.137.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.137.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.137.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.138.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.138.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.138.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.139.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.139.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.139.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.140.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.140.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.140.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.141.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.141.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.141.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.142.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.142.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.142.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.143.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.143.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.143.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.144.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.144.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.144.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.145.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.145.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.145.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.146.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.146.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.146.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.147.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.147.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.147.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.148.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.148.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.148.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.149.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.149.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.149.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.150.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.150.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.150.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.151.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.151.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.151.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.152.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.152.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.152.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.153.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.153.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.153.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.154.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.154.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.154.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.155.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.155.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.155.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.156.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.156.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.156.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.157.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.157.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.157.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.158.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.158.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.158.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.159.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.159.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.159.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.160.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.160.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.160.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.161.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.161.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.161.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.162.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.162.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.162.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.163.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.163.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.163.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.164.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.164.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.164.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.165.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.165.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.165.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.166.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.166.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.166.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.167.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.167.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.167.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.168.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.168.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.168.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.169.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.169.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.169.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.170.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.170.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.170.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.171.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.171.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.171.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.172.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.172.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.172.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.173.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.173.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.173.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.174.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.174.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.174.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.175.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.175.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.175.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.176.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.176.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.176.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.177.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.177.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.177.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.178.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.178.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.178.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.179.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.179.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.179.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.180.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.180.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.180.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.181.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.181.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.181.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.182.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.182.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.182.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.183.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.183.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.183.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.184.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.184.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.184.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.185.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.185.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.185.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.186.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.186.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.186.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.187.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.187.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.187.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.188.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.188.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.188.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.189.gate_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.189.up_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.189.down_proj.weight": "model-00102-of-000163.safetensors", + "model.layers.39.mlp.experts.190.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.190.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.190.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.191.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.191.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.191.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.192.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.192.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.192.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.193.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.193.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.193.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.194.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.194.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.194.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.195.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.195.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.195.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.196.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.196.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.196.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.197.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.197.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.197.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.198.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.198.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.198.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.199.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.199.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.199.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.200.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.200.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.200.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.201.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.201.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.201.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.202.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.202.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.202.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.203.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.203.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.203.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.204.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.204.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.204.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.205.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.205.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.205.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.206.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.206.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.206.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.207.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.207.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.207.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.208.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.208.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.208.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.209.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.209.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.209.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.210.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.210.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.210.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.211.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.211.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.211.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.212.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.212.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.212.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.213.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.213.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.213.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.214.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.214.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.214.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.215.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.215.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.215.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.216.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.216.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.216.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.217.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.217.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.217.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.218.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.218.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.218.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.219.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.219.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.219.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.220.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.220.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.220.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.221.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.221.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.221.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.222.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.222.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.222.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.223.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.223.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.223.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.224.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.224.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.224.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.225.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.225.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.225.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.226.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.226.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.226.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.227.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.227.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.227.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.228.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.228.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.228.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.229.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.229.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.229.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.230.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.230.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.230.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.231.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.231.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.231.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.232.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.232.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.232.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.233.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.233.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.233.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.234.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.234.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.234.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.235.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.235.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.235.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.236.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.236.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.236.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.237.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.237.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.237.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.238.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.238.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.238.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.239.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.239.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.239.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.240.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.240.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.240.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.241.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.241.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.241.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.242.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.242.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.242.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.243.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.243.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.243.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.244.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.244.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.244.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.245.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.245.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.245.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.246.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.246.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.246.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.247.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.247.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.247.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.248.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.248.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.248.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.249.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.249.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.249.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.250.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.250.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.250.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.251.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.251.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.251.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.252.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.252.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.252.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.253.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.253.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.253.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.254.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.254.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.254.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.255.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.255.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.mlp.experts.255.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.39.input_layernorm.weight": "model-00103-of-000163.safetensors", + "model.layers.39.post_attention_layernorm.weight": "model-00103-of-000163.safetensors", + "model.layers.40.self_attn.q_a_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.self_attn.q_a_layernorm.weight": "model-00103-of-000163.safetensors", + "model.layers.40.self_attn.q_b_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.self_attn.kv_a_proj_with_mqa.weight": "model-00103-of-000163.safetensors", + "model.layers.40.self_attn.kv_a_layernorm.weight": "model-00103-of-000163.safetensors", + "model.layers.40.self_attn.kv_b_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.self_attn.o_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.gate.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.gate.e_score_correction_bias": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.shared_experts.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.shared_experts.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.shared_experts.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.0.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.0.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.0.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.1.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.1.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.1.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.2.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.2.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.2.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.3.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.3.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.3.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.4.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.4.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.4.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.5.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.5.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.5.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.6.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.6.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.6.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.7.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.7.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.7.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.8.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.8.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.8.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.9.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.9.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.9.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.10.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.10.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.10.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.11.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.11.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.11.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.12.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.12.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.12.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.13.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.13.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.13.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.14.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.14.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.14.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.15.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.15.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.15.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.16.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.16.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.16.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.17.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.17.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.17.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.18.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.18.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.18.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.19.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.19.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.19.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.20.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.20.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.20.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.21.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.21.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.21.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.22.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.22.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.22.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.23.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.23.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.23.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.24.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.24.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.24.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.25.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.25.up_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.25.down_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.26.gate_proj.weight": "model-00103-of-000163.safetensors", + "model.layers.40.mlp.experts.26.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.26.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.27.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.27.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.27.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.28.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.28.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.28.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.29.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.29.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.29.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.30.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.30.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.30.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.31.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.31.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.31.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.32.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.32.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.32.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.33.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.33.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.33.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.34.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.34.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.34.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.35.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.35.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.35.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.36.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.36.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.36.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.37.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.37.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.37.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.38.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.38.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.38.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.39.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.39.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.39.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.40.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.40.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.40.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.41.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.41.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.41.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.42.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.42.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.42.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.43.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.43.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.43.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.44.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.44.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.44.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.45.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.45.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.45.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.46.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.46.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.46.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.47.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.47.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.47.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.48.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.48.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.48.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.49.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.49.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.49.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.50.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.50.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.50.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.51.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.51.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.51.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.52.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.52.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.52.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.53.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.53.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.53.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.54.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.54.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.54.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.55.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.55.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.55.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.56.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.56.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.56.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.57.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.57.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.57.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.58.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.58.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.58.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.59.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.59.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.59.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.60.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.60.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.60.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.61.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.61.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.61.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.62.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.62.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.62.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.63.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.63.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.63.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.64.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.64.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.64.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.65.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.65.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.65.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.66.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.66.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.66.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.67.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.67.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.67.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.68.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.68.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.68.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.69.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.69.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.69.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.70.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.70.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.70.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.71.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.71.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.71.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.72.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.72.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.72.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.73.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.73.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.73.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.74.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.74.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.74.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.75.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.75.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.75.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.76.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.76.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.76.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.77.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.77.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.77.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.78.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.78.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.78.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.79.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.79.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.79.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.80.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.80.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.80.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.81.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.81.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.81.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.82.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.82.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.82.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.83.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.83.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.83.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.84.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.84.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.84.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.85.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.85.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.85.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.86.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.86.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.86.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.87.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.87.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.87.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.88.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.88.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.88.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.89.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.89.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.89.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.90.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.90.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.90.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.91.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.91.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.91.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.92.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.92.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.92.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.93.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.93.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.93.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.94.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.94.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.94.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.95.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.95.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.95.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.96.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.96.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.96.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.97.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.97.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.97.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.98.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.98.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.98.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.99.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.99.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.99.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.100.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.100.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.100.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.101.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.101.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.101.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.102.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.102.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.102.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.103.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.103.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.103.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.104.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.104.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.104.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.105.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.105.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.105.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.106.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.106.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.106.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.107.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.107.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.107.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.108.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.108.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.108.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.109.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.109.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.109.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.110.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.110.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.110.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.111.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.111.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.111.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.112.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.112.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.112.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.113.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.113.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.113.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.114.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.114.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.114.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.115.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.115.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.115.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.116.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.116.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.116.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.117.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.117.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.117.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.118.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.118.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.118.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.119.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.119.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.119.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.120.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.120.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.120.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.121.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.121.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.121.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.122.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.122.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.122.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.123.gate_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.123.up_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.123.down_proj.weight": "model-00104-of-000163.safetensors", + "model.layers.40.mlp.experts.124.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.124.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.124.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.125.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.125.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.125.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.126.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.126.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.126.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.127.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.127.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.127.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.128.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.128.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.128.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.129.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.129.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.129.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.130.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.130.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.130.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.131.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.131.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.131.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.132.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.132.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.132.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.133.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.133.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.133.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.134.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.134.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.134.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.135.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.135.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.135.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.136.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.136.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.136.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.137.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.137.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.137.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.138.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.138.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.138.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.139.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.139.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.139.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.140.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.140.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.140.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.141.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.141.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.141.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.142.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.142.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.142.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.143.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.143.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.143.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.144.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.144.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.144.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.145.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.145.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.145.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.146.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.146.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.146.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.147.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.147.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.147.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.148.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.148.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.148.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.149.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.149.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.149.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.150.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.150.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.150.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.151.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.151.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.151.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.152.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.152.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.152.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.153.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.153.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.153.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.154.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.154.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.154.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.155.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.155.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.155.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.156.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.156.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.156.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.157.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.157.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.157.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.158.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.158.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.158.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.159.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.159.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.159.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.160.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.160.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.160.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.161.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.161.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.161.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.162.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.162.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.162.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.163.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.163.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.163.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.164.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.164.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.164.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.165.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.165.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.165.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.166.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.166.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.166.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.167.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.167.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.167.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.168.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.168.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.168.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.169.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.169.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.169.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.170.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.170.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.170.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.171.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.171.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.171.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.172.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.172.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.172.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.173.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.173.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.173.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.174.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.174.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.174.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.175.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.175.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.175.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.176.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.176.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.176.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.177.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.177.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.177.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.178.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.178.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.178.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.179.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.179.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.179.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.180.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.180.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.180.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.181.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.181.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.181.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.182.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.182.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.182.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.183.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.183.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.183.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.184.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.184.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.184.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.185.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.185.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.185.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.186.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.186.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.186.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.187.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.187.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.187.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.188.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.188.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.188.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.189.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.189.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.189.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.190.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.190.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.190.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.191.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.191.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.191.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.192.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.192.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.192.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.193.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.193.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.193.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.194.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.194.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.194.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.195.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.195.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.195.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.196.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.196.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.196.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.197.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.197.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.197.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.198.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.198.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.198.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.199.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.199.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.199.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.200.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.200.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.200.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.201.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.201.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.201.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.202.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.202.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.202.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.203.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.203.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.203.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.204.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.204.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.204.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.205.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.205.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.205.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.206.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.206.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.206.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.207.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.207.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.207.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.208.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.208.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.208.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.209.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.209.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.209.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.210.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.210.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.210.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.211.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.211.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.211.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.212.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.212.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.212.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.213.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.213.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.213.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.214.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.214.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.214.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.215.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.215.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.215.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.216.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.216.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.216.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.217.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.217.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.217.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.218.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.218.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.218.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.219.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.219.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.219.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.220.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.220.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.220.down_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.221.gate_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.221.up_proj.weight": "model-00105-of-000163.safetensors", + "model.layers.40.mlp.experts.221.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.222.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.222.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.222.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.223.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.223.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.223.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.224.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.224.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.224.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.225.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.225.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.225.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.226.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.226.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.226.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.227.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.227.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.227.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.228.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.228.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.228.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.229.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.229.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.229.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.230.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.230.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.230.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.231.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.231.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.231.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.232.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.232.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.232.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.233.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.233.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.233.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.234.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.234.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.234.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.235.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.235.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.235.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.236.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.236.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.236.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.237.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.237.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.237.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.238.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.238.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.238.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.239.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.239.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.239.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.240.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.240.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.240.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.241.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.241.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.241.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.242.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.242.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.242.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.243.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.243.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.243.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.244.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.244.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.244.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.245.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.245.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.245.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.246.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.246.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.246.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.247.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.247.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.247.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.248.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.248.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.248.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.249.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.249.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.249.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.250.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.250.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.250.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.251.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.251.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.251.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.252.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.252.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.252.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.253.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.253.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.253.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.254.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.254.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.254.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.255.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.255.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.mlp.experts.255.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.40.input_layernorm.weight": "model-00106-of-000163.safetensors", + "model.layers.40.post_attention_layernorm.weight": "model-00106-of-000163.safetensors", + "model.layers.41.self_attn.q_a_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.self_attn.q_a_layernorm.weight": "model-00106-of-000163.safetensors", + "model.layers.41.self_attn.q_b_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.self_attn.kv_a_proj_with_mqa.weight": "model-00106-of-000163.safetensors", + "model.layers.41.self_attn.kv_a_layernorm.weight": "model-00106-of-000163.safetensors", + "model.layers.41.self_attn.kv_b_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.self_attn.o_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.gate.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.gate.e_score_correction_bias": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.shared_experts.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.shared_experts.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.shared_experts.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.0.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.0.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.0.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.1.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.1.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.1.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.2.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.2.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.2.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.3.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.3.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.3.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.4.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.4.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.4.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.5.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.5.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.5.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.6.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.6.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.6.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.7.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.7.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.7.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.8.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.8.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.8.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.9.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.9.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.9.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.10.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.10.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.10.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.11.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.11.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.11.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.12.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.12.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.12.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.13.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.13.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.13.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.14.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.14.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.14.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.15.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.15.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.15.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.16.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.16.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.16.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.17.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.17.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.17.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.18.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.18.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.18.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.19.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.19.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.19.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.20.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.20.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.20.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.21.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.21.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.21.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.22.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.22.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.22.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.23.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.23.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.23.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.24.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.24.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.24.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.25.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.25.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.25.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.26.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.26.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.26.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.27.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.27.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.27.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.28.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.28.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.28.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.29.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.29.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.29.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.30.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.30.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.30.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.31.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.31.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.31.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.32.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.32.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.32.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.33.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.33.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.33.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.34.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.34.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.34.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.35.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.35.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.35.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.36.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.36.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.36.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.37.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.37.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.37.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.38.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.38.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.38.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.39.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.39.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.39.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.40.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.40.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.40.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.41.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.41.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.41.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.42.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.42.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.42.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.43.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.43.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.43.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.44.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.44.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.44.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.45.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.45.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.45.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.46.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.46.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.46.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.47.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.47.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.47.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.48.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.48.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.48.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.49.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.49.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.49.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.50.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.50.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.50.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.51.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.51.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.51.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.52.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.52.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.52.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.53.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.53.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.53.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.54.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.54.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.54.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.55.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.55.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.55.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.56.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.56.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.56.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.57.gate_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.57.up_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.57.down_proj.weight": "model-00106-of-000163.safetensors", + "model.layers.41.mlp.experts.58.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.58.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.58.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.59.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.59.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.59.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.60.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.60.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.60.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.61.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.61.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.61.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.62.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.62.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.62.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.63.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.63.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.63.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.64.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.64.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.64.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.65.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.65.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.65.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.66.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.66.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.66.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.67.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.67.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.67.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.68.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.68.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.68.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.69.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.69.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.69.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.70.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.70.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.70.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.71.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.71.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.71.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.72.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.72.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.72.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.73.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.73.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.73.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.74.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.74.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.74.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.75.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.75.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.75.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.76.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.76.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.76.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.77.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.77.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.77.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.78.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.78.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.78.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.79.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.79.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.79.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.80.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.80.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.80.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.81.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.81.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.81.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.82.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.82.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.82.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.83.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.83.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.83.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.84.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.84.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.84.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.85.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.85.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.85.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.86.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.86.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.86.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.87.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.87.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.87.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.88.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.88.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.88.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.89.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.89.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.89.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.90.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.90.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.90.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.91.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.91.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.91.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.92.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.92.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.92.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.93.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.93.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.93.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.94.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.94.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.94.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.95.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.95.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.95.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.96.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.96.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.96.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.97.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.97.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.97.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.98.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.98.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.98.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.99.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.99.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.99.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.100.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.100.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.100.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.101.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.101.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.101.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.102.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.102.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.102.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.103.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.103.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.103.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.104.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.104.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.104.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.105.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.105.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.105.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.106.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.106.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.106.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.107.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.107.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.107.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.108.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.108.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.108.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.109.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.109.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.109.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.110.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.110.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.110.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.111.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.111.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.111.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.112.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.112.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.112.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.113.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.113.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.113.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.114.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.114.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.114.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.115.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.115.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.115.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.116.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.116.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.116.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.117.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.117.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.117.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.118.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.118.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.118.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.119.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.119.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.119.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.120.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.120.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.120.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.121.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.121.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.121.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.122.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.122.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.122.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.123.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.123.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.123.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.124.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.124.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.124.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.125.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.125.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.125.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.126.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.126.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.126.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.127.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.127.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.127.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.128.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.128.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.128.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.129.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.129.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.129.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.130.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.130.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.130.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.131.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.131.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.131.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.132.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.132.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.132.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.133.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.133.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.133.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.134.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.134.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.134.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.135.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.135.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.135.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.136.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.136.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.136.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.137.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.137.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.137.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.138.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.138.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.138.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.139.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.139.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.139.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.140.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.140.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.140.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.141.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.141.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.141.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.142.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.142.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.142.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.143.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.143.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.143.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.144.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.144.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.144.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.145.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.145.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.145.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.146.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.146.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.146.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.147.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.147.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.147.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.148.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.148.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.148.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.149.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.149.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.149.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.150.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.150.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.150.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.151.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.151.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.151.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.152.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.152.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.152.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.153.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.153.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.153.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.154.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.154.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.154.down_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.155.gate_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.155.up_proj.weight": "model-00107-of-000163.safetensors", + "model.layers.41.mlp.experts.155.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.156.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.156.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.156.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.157.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.157.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.157.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.158.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.158.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.158.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.159.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.159.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.159.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.160.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.160.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.160.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.161.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.161.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.161.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.162.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.162.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.162.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.163.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.163.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.163.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.164.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.164.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.164.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.165.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.165.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.165.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.166.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.166.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.166.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.167.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.167.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.167.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.168.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.168.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.168.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.169.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.169.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.169.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.170.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.170.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.170.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.171.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.171.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.171.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.172.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.172.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.172.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.173.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.173.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.173.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.174.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.174.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.174.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.175.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.175.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.175.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.176.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.176.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.176.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.177.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.177.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.177.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.178.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.178.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.178.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.179.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.179.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.179.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.180.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.180.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.180.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.181.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.181.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.181.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.182.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.182.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.182.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.183.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.183.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.183.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.184.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.184.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.184.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.185.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.185.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.185.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.186.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.186.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.186.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.187.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.187.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.187.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.188.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.188.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.188.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.189.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.189.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.189.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.190.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.190.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.190.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.191.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.191.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.191.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.192.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.192.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.192.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.193.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.193.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.193.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.194.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.194.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.194.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.195.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.195.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.195.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.196.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.196.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.196.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.197.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.197.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.197.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.198.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.198.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.198.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.199.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.199.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.199.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.200.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.200.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.200.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.201.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.201.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.201.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.202.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.202.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.202.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.203.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.203.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.203.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.204.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.204.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.204.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.205.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.205.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.205.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.206.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.206.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.206.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.207.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.207.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.207.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.208.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.208.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.208.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.209.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.209.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.209.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.210.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.210.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.210.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.211.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.211.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.211.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.212.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.212.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.212.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.213.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.213.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.213.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.214.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.214.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.214.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.215.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.215.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.215.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.216.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.216.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.216.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.217.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.217.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.217.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.218.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.218.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.218.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.219.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.219.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.219.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.220.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.220.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.220.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.221.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.221.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.221.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.222.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.222.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.222.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.223.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.223.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.223.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.224.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.224.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.224.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.225.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.225.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.225.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.226.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.226.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.226.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.227.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.227.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.227.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.228.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.228.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.228.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.229.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.229.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.229.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.230.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.230.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.230.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.231.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.231.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.231.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.232.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.232.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.232.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.233.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.233.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.233.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.234.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.234.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.234.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.235.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.235.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.235.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.236.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.236.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.236.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.237.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.237.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.237.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.238.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.238.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.238.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.239.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.239.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.239.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.240.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.240.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.240.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.241.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.241.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.241.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.242.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.242.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.242.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.243.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.243.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.243.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.244.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.244.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.244.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.245.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.245.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.245.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.246.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.246.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.246.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.247.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.247.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.247.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.248.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.248.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.248.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.249.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.249.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.249.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.250.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.250.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.250.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.251.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.251.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.251.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.252.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.252.up_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.252.down_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.253.gate_proj.weight": "model-00108-of-000163.safetensors", + "model.layers.41.mlp.experts.253.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.41.mlp.experts.253.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.41.mlp.experts.254.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.41.mlp.experts.254.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.41.mlp.experts.254.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.41.mlp.experts.255.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.41.mlp.experts.255.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.41.mlp.experts.255.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.41.input_layernorm.weight": "model-00109-of-000163.safetensors", + "model.layers.41.post_attention_layernorm.weight": "model-00109-of-000163.safetensors", + "model.layers.42.self_attn.q_a_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.self_attn.q_a_layernorm.weight": "model-00109-of-000163.safetensors", + "model.layers.42.self_attn.q_b_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.self_attn.kv_a_proj_with_mqa.weight": "model-00109-of-000163.safetensors", + "model.layers.42.self_attn.kv_a_layernorm.weight": "model-00109-of-000163.safetensors", + "model.layers.42.self_attn.kv_b_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.self_attn.o_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.gate.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.gate.e_score_correction_bias": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.shared_experts.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.shared_experts.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.shared_experts.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.0.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.0.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.0.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.1.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.1.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.1.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.2.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.2.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.2.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.3.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.3.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.3.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.4.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.4.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.4.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.5.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.5.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.5.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.6.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.6.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.6.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.7.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.7.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.7.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.8.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.8.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.8.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.9.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.9.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.9.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.10.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.10.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.10.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.11.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.11.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.11.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.12.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.12.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.12.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.13.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.13.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.13.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.14.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.14.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.14.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.15.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.15.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.15.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.16.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.16.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.16.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.17.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.17.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.17.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.18.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.18.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.18.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.19.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.19.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.19.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.20.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.20.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.20.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.21.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.21.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.21.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.22.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.22.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.22.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.23.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.23.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.23.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.24.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.24.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.24.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.25.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.25.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.25.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.26.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.26.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.26.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.27.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.27.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.27.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.28.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.28.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.28.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.29.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.29.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.29.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.30.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.30.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.30.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.31.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.31.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.31.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.32.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.32.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.32.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.33.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.33.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.33.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.34.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.34.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.34.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.35.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.35.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.35.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.36.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.36.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.36.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.37.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.37.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.37.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.38.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.38.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.38.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.39.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.39.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.39.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.40.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.40.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.40.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.41.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.41.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.41.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.42.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.42.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.42.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.43.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.43.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.43.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.44.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.44.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.44.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.45.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.45.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.45.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.46.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.46.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.46.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.47.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.47.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.47.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.48.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.48.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.48.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.49.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.49.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.49.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.50.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.50.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.50.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.51.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.51.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.51.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.52.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.52.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.52.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.53.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.53.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.53.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.54.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.54.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.54.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.55.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.55.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.55.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.56.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.56.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.56.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.57.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.57.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.57.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.58.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.58.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.58.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.59.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.59.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.59.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.60.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.60.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.60.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.61.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.61.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.61.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.62.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.62.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.62.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.63.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.63.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.63.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.64.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.64.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.64.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.65.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.65.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.65.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.66.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.66.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.66.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.67.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.67.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.67.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.68.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.68.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.68.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.69.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.69.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.69.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.70.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.70.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.70.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.71.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.71.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.71.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.72.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.72.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.72.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.73.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.73.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.73.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.74.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.74.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.74.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.75.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.75.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.75.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.76.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.76.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.76.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.77.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.77.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.77.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.78.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.78.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.78.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.79.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.79.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.79.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.80.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.80.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.80.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.81.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.81.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.81.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.82.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.82.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.82.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.83.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.83.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.83.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.84.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.84.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.84.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.85.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.85.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.85.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.86.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.86.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.86.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.87.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.87.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.87.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.88.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.88.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.88.down_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.89.gate_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.89.up_proj.weight": "model-00109-of-000163.safetensors", + "model.layers.42.mlp.experts.89.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.90.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.90.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.90.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.91.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.91.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.91.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.92.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.92.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.92.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.93.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.93.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.93.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.94.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.94.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.94.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.95.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.95.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.95.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.96.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.96.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.96.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.97.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.97.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.97.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.98.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.98.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.98.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.99.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.99.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.99.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.100.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.100.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.100.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.101.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.101.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.101.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.102.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.102.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.102.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.103.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.103.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.103.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.104.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.104.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.104.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.105.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.105.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.105.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.106.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.106.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.106.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.107.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.107.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.107.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.108.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.108.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.108.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.109.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.109.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.109.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.110.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.110.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.110.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.111.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.111.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.111.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.112.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.112.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.112.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.113.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.113.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.113.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.114.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.114.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.114.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.115.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.115.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.115.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.116.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.116.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.116.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.117.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.117.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.117.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.118.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.118.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.118.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.119.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.119.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.119.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.120.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.120.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.120.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.121.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.121.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.121.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.122.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.122.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.122.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.123.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.123.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.123.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.124.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.124.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.124.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.125.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.125.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.125.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.126.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.126.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.126.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.127.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.127.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.127.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.128.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.128.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.128.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.129.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.129.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.129.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.130.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.130.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.130.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.131.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.131.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.131.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.132.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.132.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.132.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.133.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.133.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.133.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.134.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.134.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.134.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.135.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.135.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.135.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.136.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.136.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.136.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.137.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.137.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.137.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.138.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.138.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.138.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.139.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.139.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.139.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.140.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.140.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.140.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.141.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.141.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.141.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.142.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.142.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.142.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.143.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.143.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.143.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.144.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.144.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.144.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.145.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.145.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.145.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.146.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.146.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.146.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.147.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.147.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.147.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.148.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.148.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.148.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.149.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.149.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.149.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.150.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.150.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.150.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.151.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.151.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.151.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.152.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.152.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.152.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.153.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.153.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.153.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.154.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.154.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.154.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.155.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.155.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.155.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.156.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.156.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.156.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.157.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.157.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.157.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.158.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.158.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.158.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.159.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.159.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.159.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.160.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.160.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.160.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.161.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.161.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.161.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.162.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.162.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.162.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.163.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.163.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.163.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.164.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.164.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.164.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.165.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.165.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.165.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.166.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.166.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.166.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.167.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.167.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.167.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.168.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.168.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.168.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.169.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.169.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.169.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.170.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.170.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.170.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.171.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.171.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.171.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.172.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.172.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.172.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.173.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.173.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.173.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.174.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.174.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.174.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.175.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.175.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.175.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.176.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.176.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.176.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.177.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.177.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.177.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.178.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.178.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.178.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.179.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.179.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.179.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.180.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.180.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.180.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.181.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.181.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.181.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.182.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.182.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.182.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.183.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.183.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.183.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.184.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.184.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.184.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.185.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.185.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.185.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.186.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.186.up_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.186.down_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.187.gate_proj.weight": "model-00110-of-000163.safetensors", + "model.layers.42.mlp.experts.187.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.187.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.188.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.188.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.188.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.189.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.189.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.189.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.190.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.190.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.190.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.191.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.191.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.191.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.192.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.192.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.192.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.193.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.193.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.193.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.194.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.194.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.194.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.195.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.195.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.195.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.196.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.196.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.196.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.197.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.197.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.197.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.198.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.198.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.198.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.199.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.199.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.199.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.200.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.200.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.200.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.201.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.201.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.201.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.202.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.202.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.202.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.203.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.203.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.203.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.204.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.204.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.204.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.205.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.205.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.205.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.206.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.206.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.206.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.207.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.207.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.207.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.208.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.208.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.208.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.209.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.209.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.209.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.210.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.210.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.210.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.211.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.211.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.211.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.212.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.212.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.212.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.213.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.213.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.213.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.214.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.214.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.214.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.215.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.215.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.215.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.216.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.216.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.216.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.217.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.217.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.217.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.218.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.218.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.218.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.219.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.219.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.219.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.220.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.220.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.220.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.221.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.221.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.221.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.222.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.222.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.222.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.223.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.223.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.223.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.224.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.224.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.224.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.225.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.225.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.225.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.226.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.226.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.226.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.227.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.227.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.227.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.228.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.228.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.228.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.229.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.229.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.229.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.230.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.230.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.230.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.231.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.231.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.231.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.232.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.232.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.232.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.233.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.233.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.233.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.234.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.234.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.234.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.235.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.235.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.235.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.236.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.236.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.236.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.237.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.237.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.237.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.238.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.238.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.238.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.239.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.239.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.239.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.240.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.240.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.240.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.241.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.241.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.241.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.242.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.242.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.242.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.243.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.243.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.243.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.244.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.244.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.244.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.245.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.245.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.245.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.246.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.246.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.246.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.247.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.247.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.247.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.248.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.248.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.248.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.249.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.249.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.249.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.250.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.250.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.250.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.251.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.251.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.251.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.252.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.252.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.252.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.253.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.253.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.253.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.254.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.254.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.254.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.255.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.255.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.mlp.experts.255.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.42.input_layernorm.weight": "model-00111-of-000163.safetensors", + "model.layers.42.post_attention_layernorm.weight": "model-00111-of-000163.safetensors", + "model.layers.43.self_attn.q_a_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.self_attn.q_a_layernorm.weight": "model-00111-of-000163.safetensors", + "model.layers.43.self_attn.q_b_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.self_attn.kv_a_proj_with_mqa.weight": "model-00111-of-000163.safetensors", + "model.layers.43.self_attn.kv_a_layernorm.weight": "model-00111-of-000163.safetensors", + "model.layers.43.self_attn.kv_b_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.self_attn.o_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.gate.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.gate.e_score_correction_bias": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.shared_experts.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.shared_experts.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.shared_experts.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.0.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.0.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.0.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.1.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.1.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.1.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.2.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.2.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.2.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.3.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.3.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.3.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.4.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.4.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.4.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.5.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.5.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.5.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.6.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.6.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.6.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.7.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.7.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.7.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.8.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.8.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.8.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.9.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.9.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.9.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.10.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.10.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.10.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.11.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.11.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.11.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.12.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.12.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.12.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.13.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.13.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.13.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.14.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.14.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.14.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.15.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.15.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.15.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.16.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.16.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.16.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.17.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.17.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.17.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.18.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.18.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.18.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.19.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.19.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.19.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.20.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.20.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.20.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.21.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.21.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.21.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.22.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.22.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.22.down_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.23.gate_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.23.up_proj.weight": "model-00111-of-000163.safetensors", + "model.layers.43.mlp.experts.23.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.24.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.24.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.24.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.25.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.25.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.25.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.26.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.26.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.26.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.27.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.27.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.27.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.28.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.28.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.28.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.29.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.29.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.29.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.30.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.30.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.30.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.31.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.31.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.31.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.32.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.32.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.32.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.33.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.33.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.33.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.34.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.34.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.34.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.35.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.35.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.35.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.36.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.36.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.36.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.37.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.37.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.37.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.38.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.38.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.38.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.39.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.39.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.39.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.40.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.40.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.40.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.41.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.41.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.41.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.42.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.42.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.42.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.43.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.43.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.43.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.44.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.44.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.44.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.45.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.45.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.45.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.46.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.46.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.46.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.47.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.47.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.47.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.48.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.48.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.48.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.49.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.49.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.49.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.50.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.50.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.50.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.51.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.51.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.51.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.52.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.52.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.52.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.53.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.53.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.53.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.54.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.54.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.54.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.55.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.55.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.55.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.56.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.56.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.56.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.57.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.57.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.57.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.58.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.58.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.58.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.59.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.59.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.59.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.60.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.60.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.60.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.61.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.61.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.61.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.62.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.62.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.62.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.63.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.63.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.63.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.64.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.64.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.64.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.65.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.65.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.65.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.66.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.66.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.66.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.67.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.67.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.67.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.68.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.68.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.68.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.69.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.69.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.69.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.70.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.70.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.70.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.71.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.71.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.71.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.72.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.72.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.72.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.73.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.73.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.73.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.74.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.74.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.74.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.75.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.75.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.75.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.76.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.76.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.76.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.77.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.77.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.77.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.78.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.78.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.78.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.79.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.79.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.79.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.80.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.80.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.80.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.81.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.81.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.81.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.82.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.82.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.82.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.83.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.83.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.83.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.84.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.84.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.84.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.85.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.85.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.85.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.86.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.86.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.86.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.87.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.87.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.87.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.88.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.88.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.88.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.89.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.89.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.89.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.90.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.90.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.90.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.91.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.91.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.91.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.92.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.92.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.92.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.93.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.93.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.93.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.94.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.94.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.94.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.95.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.95.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.95.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.96.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.96.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.96.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.97.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.97.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.97.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.98.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.98.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.98.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.99.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.99.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.99.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.100.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.100.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.100.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.101.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.101.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.101.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.102.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.102.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.102.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.103.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.103.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.103.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.104.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.104.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.104.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.105.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.105.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.105.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.106.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.106.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.106.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.107.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.107.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.107.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.108.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.108.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.108.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.109.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.109.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.109.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.110.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.110.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.110.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.111.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.111.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.111.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.112.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.112.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.112.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.113.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.113.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.113.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.114.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.114.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.114.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.115.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.115.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.115.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.116.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.116.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.116.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.117.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.117.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.117.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.118.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.118.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.118.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.119.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.119.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.119.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.120.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.120.up_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.120.down_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.121.gate_proj.weight": "model-00112-of-000163.safetensors", + "model.layers.43.mlp.experts.121.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.121.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.122.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.122.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.122.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.123.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.123.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.123.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.124.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.124.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.124.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.125.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.125.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.125.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.126.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.126.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.126.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.127.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.127.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.127.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.128.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.128.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.128.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.129.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.129.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.129.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.130.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.130.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.130.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.131.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.131.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.131.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.132.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.132.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.132.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.133.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.133.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.133.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.134.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.134.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.134.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.135.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.135.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.135.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.136.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.136.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.136.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.137.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.137.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.137.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.138.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.138.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.138.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.139.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.139.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.139.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.140.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.140.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.140.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.141.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.141.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.141.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.142.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.142.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.142.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.143.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.143.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.143.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.144.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.144.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.144.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.145.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.145.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.145.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.146.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.146.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.146.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.147.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.147.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.147.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.148.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.148.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.148.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.149.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.149.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.149.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.150.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.150.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.150.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.151.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.151.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.151.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.152.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.152.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.152.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.153.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.153.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.153.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.154.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.154.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.154.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.155.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.155.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.155.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.156.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.156.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.156.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.157.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.157.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.157.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.158.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.158.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.158.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.159.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.159.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.159.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.160.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.160.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.160.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.161.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.161.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.161.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.162.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.162.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.162.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.163.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.163.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.163.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.164.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.164.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.164.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.165.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.165.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.165.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.166.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.166.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.166.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.167.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.167.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.167.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.168.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.168.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.168.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.169.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.169.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.169.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.170.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.170.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.170.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.171.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.171.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.171.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.172.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.172.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.172.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.173.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.173.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.173.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.174.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.174.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.174.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.175.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.175.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.175.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.176.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.176.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.176.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.177.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.177.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.177.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.178.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.178.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.178.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.179.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.179.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.179.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.180.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.180.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.180.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.181.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.181.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.181.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.182.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.182.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.182.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.183.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.183.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.183.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.184.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.184.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.184.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.185.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.185.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.185.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.186.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.186.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.186.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.187.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.187.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.187.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.188.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.188.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.188.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.189.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.189.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.189.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.190.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.190.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.190.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.191.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.191.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.191.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.192.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.192.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.192.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.193.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.193.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.193.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.194.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.194.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.194.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.195.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.195.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.195.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.196.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.196.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.196.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.197.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.197.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.197.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.198.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.198.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.198.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.199.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.199.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.199.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.200.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.200.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.200.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.201.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.201.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.201.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.202.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.202.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.202.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.203.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.203.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.203.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.204.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.204.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.204.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.205.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.205.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.205.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.206.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.206.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.206.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.207.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.207.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.207.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.208.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.208.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.208.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.209.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.209.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.209.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.210.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.210.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.210.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.211.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.211.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.211.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.212.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.212.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.212.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.213.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.213.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.213.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.214.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.214.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.214.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.215.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.215.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.215.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.216.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.216.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.216.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.217.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.217.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.217.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.218.gate_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.218.up_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.218.down_proj.weight": "model-00113-of-000163.safetensors", + "model.layers.43.mlp.experts.219.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.219.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.219.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.220.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.220.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.220.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.221.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.221.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.221.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.222.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.222.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.222.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.223.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.223.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.223.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.224.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.224.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.224.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.225.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.225.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.225.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.226.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.226.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.226.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.227.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.227.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.227.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.228.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.228.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.228.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.229.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.229.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.229.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.230.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.230.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.230.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.231.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.231.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.231.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.232.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.232.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.232.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.233.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.233.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.233.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.234.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.234.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.234.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.235.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.235.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.235.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.236.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.236.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.236.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.237.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.237.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.237.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.238.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.238.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.238.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.239.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.239.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.239.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.240.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.240.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.240.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.241.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.241.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.241.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.242.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.242.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.242.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.243.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.243.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.243.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.244.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.244.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.244.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.245.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.245.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.245.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.246.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.246.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.246.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.247.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.247.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.247.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.248.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.248.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.248.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.249.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.249.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.249.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.250.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.250.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.250.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.251.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.251.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.251.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.252.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.252.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.252.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.253.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.253.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.253.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.254.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.254.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.254.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.255.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.255.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.mlp.experts.255.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.43.input_layernorm.weight": "model-00114-of-000163.safetensors", + "model.layers.43.post_attention_layernorm.weight": "model-00114-of-000163.safetensors", + "model.layers.44.self_attn.q_a_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.self_attn.q_a_layernorm.weight": "model-00114-of-000163.safetensors", + "model.layers.44.self_attn.q_b_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.self_attn.kv_a_proj_with_mqa.weight": "model-00114-of-000163.safetensors", + "model.layers.44.self_attn.kv_a_layernorm.weight": "model-00114-of-000163.safetensors", + "model.layers.44.self_attn.kv_b_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.self_attn.o_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.gate.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.gate.e_score_correction_bias": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.shared_experts.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.shared_experts.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.shared_experts.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.0.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.0.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.0.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.1.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.1.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.1.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.2.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.2.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.2.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.3.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.3.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.3.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.4.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.4.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.4.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.5.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.5.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.5.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.6.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.6.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.6.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.7.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.7.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.7.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.8.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.8.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.8.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.9.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.9.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.9.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.10.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.10.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.10.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.11.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.11.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.11.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.12.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.12.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.12.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.13.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.13.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.13.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.14.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.14.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.14.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.15.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.15.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.15.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.16.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.16.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.16.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.17.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.17.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.17.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.18.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.18.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.18.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.19.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.19.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.19.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.20.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.20.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.20.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.21.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.21.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.21.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.22.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.22.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.22.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.23.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.23.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.23.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.24.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.24.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.24.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.25.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.25.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.25.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.26.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.26.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.26.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.27.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.27.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.27.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.28.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.28.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.28.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.29.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.29.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.29.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.30.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.30.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.30.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.31.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.31.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.31.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.32.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.32.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.32.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.33.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.33.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.33.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.34.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.34.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.34.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.35.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.35.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.35.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.36.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.36.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.36.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.37.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.37.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.37.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.38.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.38.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.38.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.39.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.39.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.39.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.40.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.40.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.40.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.41.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.41.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.41.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.42.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.42.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.42.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.43.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.43.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.43.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.44.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.44.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.44.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.45.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.45.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.45.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.46.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.46.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.46.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.47.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.47.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.47.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.48.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.48.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.48.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.49.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.49.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.49.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.50.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.50.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.50.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.51.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.51.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.51.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.52.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.52.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.52.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.53.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.53.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.53.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.54.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.54.up_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.54.down_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.55.gate_proj.weight": "model-00114-of-000163.safetensors", + "model.layers.44.mlp.experts.55.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.55.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.56.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.56.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.56.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.57.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.57.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.57.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.58.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.58.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.58.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.59.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.59.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.59.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.60.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.60.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.60.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.61.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.61.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.61.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.62.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.62.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.62.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.63.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.63.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.63.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.64.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.64.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.64.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.65.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.65.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.65.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.66.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.66.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.66.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.67.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.67.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.67.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.68.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.68.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.68.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.69.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.69.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.69.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.70.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.70.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.70.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.71.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.71.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.71.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.72.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.72.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.72.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.73.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.73.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.73.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.74.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.74.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.74.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.75.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.75.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.75.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.76.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.76.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.76.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.77.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.77.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.77.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.78.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.78.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.78.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.79.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.79.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.79.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.80.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.80.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.80.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.81.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.81.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.81.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.82.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.82.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.82.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.83.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.83.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.83.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.84.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.84.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.84.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.85.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.85.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.85.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.86.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.86.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.86.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.87.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.87.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.87.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.88.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.88.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.88.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.89.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.89.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.89.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.90.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.90.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.90.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.91.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.91.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.91.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.92.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.92.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.92.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.93.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.93.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.93.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.94.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.94.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.94.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.95.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.95.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.95.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.96.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.96.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.96.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.97.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.97.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.97.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.98.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.98.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.98.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.99.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.99.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.99.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.100.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.100.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.100.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.101.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.101.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.101.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.102.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.102.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.102.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.103.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.103.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.103.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.104.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.104.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.104.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.105.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.105.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.105.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.106.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.106.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.106.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.107.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.107.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.107.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.108.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.108.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.108.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.109.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.109.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.109.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.110.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.110.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.110.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.111.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.111.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.111.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.112.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.112.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.112.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.113.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.113.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.113.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.114.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.114.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.114.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.115.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.115.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.115.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.116.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.116.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.116.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.117.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.117.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.117.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.118.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.118.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.118.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.119.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.119.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.119.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.120.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.120.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.120.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.121.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.121.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.121.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.122.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.122.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.122.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.123.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.123.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.123.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.124.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.124.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.124.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.125.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.125.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.125.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.126.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.126.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.126.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.127.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.127.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.127.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.128.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.128.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.128.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.129.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.129.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.129.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.130.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.130.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.130.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.131.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.131.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.131.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.132.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.132.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.132.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.133.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.133.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.133.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.134.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.134.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.134.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.135.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.135.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.135.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.136.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.136.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.136.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.137.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.137.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.137.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.138.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.138.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.138.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.139.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.139.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.139.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.140.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.140.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.140.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.141.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.141.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.141.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.142.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.142.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.142.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.143.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.143.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.143.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.144.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.144.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.144.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.145.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.145.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.145.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.146.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.146.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.146.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.147.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.147.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.147.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.148.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.148.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.148.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.149.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.149.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.149.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.150.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.150.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.150.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.151.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.151.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.151.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.152.gate_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.152.up_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.152.down_proj.weight": "model-00115-of-000163.safetensors", + "model.layers.44.mlp.experts.153.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.153.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.153.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.154.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.154.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.154.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.155.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.155.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.155.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.156.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.156.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.156.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.157.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.157.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.157.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.158.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.158.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.158.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.159.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.159.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.159.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.160.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.160.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.160.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.161.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.161.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.161.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.162.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.162.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.162.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.163.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.163.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.163.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.164.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.164.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.164.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.165.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.165.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.165.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.166.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.166.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.166.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.167.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.167.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.167.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.168.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.168.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.168.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.169.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.169.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.169.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.170.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.170.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.170.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.171.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.171.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.171.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.172.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.172.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.172.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.173.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.173.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.173.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.174.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.174.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.174.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.175.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.175.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.175.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.176.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.176.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.176.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.177.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.177.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.177.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.178.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.178.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.178.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.179.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.179.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.179.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.180.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.180.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.180.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.181.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.181.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.181.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.182.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.182.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.182.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.183.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.183.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.183.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.184.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.184.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.184.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.185.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.185.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.185.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.186.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.186.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.186.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.187.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.187.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.187.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.188.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.188.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.188.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.189.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.189.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.189.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.190.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.190.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.190.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.191.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.191.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.191.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.192.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.192.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.192.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.193.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.193.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.193.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.194.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.194.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.194.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.195.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.195.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.195.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.196.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.196.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.196.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.197.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.197.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.197.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.198.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.198.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.198.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.199.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.199.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.199.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.200.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.200.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.200.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.201.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.201.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.201.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.202.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.202.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.202.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.203.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.203.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.203.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.204.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.204.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.204.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.205.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.205.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.205.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.206.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.206.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.206.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.207.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.207.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.207.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.208.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.208.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.208.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.209.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.209.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.209.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.210.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.210.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.210.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.211.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.211.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.211.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.212.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.212.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.212.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.213.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.213.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.213.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.214.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.214.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.214.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.215.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.215.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.215.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.216.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.216.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.216.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.217.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.217.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.217.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.218.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.218.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.218.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.219.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.219.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.219.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.220.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.220.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.220.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.221.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.221.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.221.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.222.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.222.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.222.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.223.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.223.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.223.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.224.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.224.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.224.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.225.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.225.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.225.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.226.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.226.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.226.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.227.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.227.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.227.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.228.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.228.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.228.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.229.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.229.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.229.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.230.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.230.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.230.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.231.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.231.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.231.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.232.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.232.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.232.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.233.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.233.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.233.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.234.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.234.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.234.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.235.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.235.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.235.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.236.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.236.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.236.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.237.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.237.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.237.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.238.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.238.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.238.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.239.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.239.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.239.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.240.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.240.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.240.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.241.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.241.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.241.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.242.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.242.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.242.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.243.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.243.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.243.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.244.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.244.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.244.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.245.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.245.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.245.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.246.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.246.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.246.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.247.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.247.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.247.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.248.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.248.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.248.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.249.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.249.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.249.down_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.250.gate_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.250.up_proj.weight": "model-00116-of-000163.safetensors", + "model.layers.44.mlp.experts.250.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.44.mlp.experts.251.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.44.mlp.experts.251.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.44.mlp.experts.251.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.44.mlp.experts.252.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.44.mlp.experts.252.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.44.mlp.experts.252.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.44.mlp.experts.253.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.44.mlp.experts.253.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.44.mlp.experts.253.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.44.mlp.experts.254.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.44.mlp.experts.254.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.44.mlp.experts.254.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.44.mlp.experts.255.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.44.mlp.experts.255.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.44.mlp.experts.255.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.44.input_layernorm.weight": "model-00117-of-000163.safetensors", + "model.layers.44.post_attention_layernorm.weight": "model-00117-of-000163.safetensors", + "model.layers.45.self_attn.q_a_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.self_attn.q_a_layernorm.weight": "model-00117-of-000163.safetensors", + "model.layers.45.self_attn.q_b_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.self_attn.kv_a_proj_with_mqa.weight": "model-00117-of-000163.safetensors", + "model.layers.45.self_attn.kv_a_layernorm.weight": "model-00117-of-000163.safetensors", + "model.layers.45.self_attn.kv_b_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.self_attn.o_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.gate.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.gate.e_score_correction_bias": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.shared_experts.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.shared_experts.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.shared_experts.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.0.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.0.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.0.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.1.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.1.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.1.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.2.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.2.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.2.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.3.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.3.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.3.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.4.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.4.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.4.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.5.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.5.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.5.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.6.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.6.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.6.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.7.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.7.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.7.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.8.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.8.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.8.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.9.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.9.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.9.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.10.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.10.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.10.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.11.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.11.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.11.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.12.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.12.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.12.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.13.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.13.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.13.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.14.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.14.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.14.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.15.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.15.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.15.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.16.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.16.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.16.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.17.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.17.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.17.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.18.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.18.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.18.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.19.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.19.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.19.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.20.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.20.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.20.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.21.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.21.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.21.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.22.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.22.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.22.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.23.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.23.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.23.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.24.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.24.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.24.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.25.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.25.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.25.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.26.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.26.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.26.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.27.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.27.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.27.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.28.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.28.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.28.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.29.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.29.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.29.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.30.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.30.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.30.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.31.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.31.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.31.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.32.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.32.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.32.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.33.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.33.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.33.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.34.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.34.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.34.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.35.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.35.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.35.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.36.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.36.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.36.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.37.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.37.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.37.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.38.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.38.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.38.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.39.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.39.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.39.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.40.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.40.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.40.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.41.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.41.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.41.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.42.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.42.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.42.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.43.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.43.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.43.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.44.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.44.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.44.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.45.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.45.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.45.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.46.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.46.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.46.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.47.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.47.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.47.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.48.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.48.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.48.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.49.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.49.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.49.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.50.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.50.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.50.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.51.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.51.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.51.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.52.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.52.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.52.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.53.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.53.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.53.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.54.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.54.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.54.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.55.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.55.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.55.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.56.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.56.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.56.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.57.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.57.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.57.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.58.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.58.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.58.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.59.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.59.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.59.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.60.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.60.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.60.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.61.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.61.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.61.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.62.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.62.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.62.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.63.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.63.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.63.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.64.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.64.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.64.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.65.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.65.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.65.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.66.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.66.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.66.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.67.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.67.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.67.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.68.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.68.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.68.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.69.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.69.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.69.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.70.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.70.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.70.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.71.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.71.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.71.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.72.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.72.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.72.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.73.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.73.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.73.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.74.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.74.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.74.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.75.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.75.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.75.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.76.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.76.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.76.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.77.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.77.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.77.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.78.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.78.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.78.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.79.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.79.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.79.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.80.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.80.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.80.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.81.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.81.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.81.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.82.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.82.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.82.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.83.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.83.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.83.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.84.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.84.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.84.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.85.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.85.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.85.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.86.gate_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.86.up_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.86.down_proj.weight": "model-00117-of-000163.safetensors", + "model.layers.45.mlp.experts.87.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.87.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.87.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.88.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.88.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.88.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.89.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.89.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.89.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.90.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.90.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.90.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.91.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.91.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.91.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.92.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.92.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.92.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.93.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.93.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.93.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.94.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.94.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.94.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.95.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.95.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.95.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.96.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.96.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.96.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.97.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.97.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.97.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.98.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.98.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.98.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.99.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.99.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.99.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.100.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.100.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.100.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.101.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.101.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.101.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.102.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.102.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.102.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.103.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.103.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.103.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.104.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.104.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.104.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.105.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.105.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.105.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.106.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.106.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.106.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.107.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.107.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.107.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.108.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.108.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.108.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.109.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.109.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.109.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.110.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.110.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.110.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.111.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.111.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.111.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.112.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.112.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.112.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.113.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.113.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.113.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.114.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.114.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.114.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.115.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.115.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.115.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.116.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.116.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.116.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.117.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.117.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.117.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.118.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.118.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.118.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.119.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.119.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.119.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.120.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.120.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.120.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.121.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.121.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.121.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.122.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.122.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.122.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.123.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.123.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.123.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.124.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.124.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.124.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.125.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.125.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.125.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.126.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.126.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.126.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.127.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.127.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.127.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.128.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.128.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.128.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.129.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.129.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.129.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.130.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.130.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.130.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.131.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.131.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.131.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.132.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.132.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.132.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.133.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.133.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.133.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.134.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.134.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.134.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.135.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.135.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.135.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.136.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.136.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.136.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.137.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.137.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.137.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.138.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.138.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.138.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.139.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.139.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.139.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.140.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.140.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.140.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.141.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.141.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.141.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.142.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.142.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.142.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.143.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.143.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.143.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.144.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.144.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.144.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.145.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.145.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.145.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.146.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.146.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.146.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.147.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.147.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.147.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.148.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.148.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.148.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.149.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.149.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.149.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.150.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.150.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.150.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.151.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.151.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.151.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.152.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.152.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.152.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.153.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.153.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.153.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.154.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.154.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.154.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.155.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.155.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.155.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.156.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.156.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.156.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.157.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.157.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.157.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.158.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.158.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.158.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.159.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.159.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.159.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.160.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.160.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.160.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.161.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.161.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.161.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.162.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.162.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.162.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.163.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.163.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.163.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.164.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.164.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.164.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.165.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.165.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.165.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.166.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.166.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.166.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.167.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.167.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.167.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.168.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.168.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.168.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.169.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.169.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.169.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.170.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.170.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.170.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.171.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.171.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.171.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.172.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.172.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.172.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.173.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.173.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.173.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.174.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.174.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.174.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.175.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.175.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.175.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.176.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.176.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.176.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.177.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.177.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.177.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.178.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.178.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.178.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.179.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.179.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.179.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.180.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.180.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.180.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.181.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.181.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.181.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.182.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.182.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.182.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.183.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.183.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.183.down_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.184.gate_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.184.up_proj.weight": "model-00118-of-000163.safetensors", + "model.layers.45.mlp.experts.184.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.185.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.185.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.185.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.186.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.186.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.186.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.187.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.187.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.187.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.188.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.188.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.188.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.189.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.189.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.189.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.190.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.190.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.190.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.191.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.191.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.191.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.192.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.192.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.192.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.193.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.193.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.193.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.194.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.194.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.194.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.195.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.195.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.195.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.196.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.196.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.196.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.197.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.197.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.197.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.198.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.198.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.198.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.199.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.199.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.199.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.200.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.200.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.200.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.201.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.201.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.201.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.202.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.202.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.202.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.203.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.203.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.203.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.204.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.204.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.204.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.205.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.205.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.205.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.206.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.206.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.206.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.207.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.207.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.207.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.208.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.208.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.208.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.209.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.209.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.209.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.210.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.210.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.210.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.211.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.211.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.211.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.212.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.212.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.212.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.213.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.213.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.213.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.214.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.214.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.214.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.215.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.215.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.215.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.216.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.216.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.216.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.217.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.217.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.217.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.218.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.218.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.218.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.219.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.219.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.219.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.220.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.220.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.220.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.221.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.221.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.221.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.222.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.222.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.222.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.223.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.223.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.223.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.224.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.224.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.224.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.225.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.225.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.225.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.226.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.226.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.226.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.227.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.227.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.227.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.228.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.228.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.228.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.229.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.229.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.229.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.230.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.230.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.230.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.231.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.231.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.231.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.232.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.232.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.232.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.233.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.233.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.233.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.234.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.234.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.234.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.235.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.235.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.235.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.236.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.236.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.236.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.237.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.237.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.237.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.238.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.238.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.238.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.239.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.239.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.239.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.240.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.240.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.240.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.241.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.241.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.241.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.242.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.242.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.242.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.243.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.243.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.243.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.244.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.244.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.244.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.245.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.245.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.245.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.246.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.246.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.246.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.247.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.247.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.247.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.248.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.248.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.248.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.249.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.249.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.249.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.250.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.250.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.250.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.251.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.251.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.251.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.252.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.252.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.252.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.253.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.253.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.253.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.254.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.254.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.254.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.255.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.255.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.mlp.experts.255.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.45.input_layernorm.weight": "model-00119-of-000163.safetensors", + "model.layers.45.post_attention_layernorm.weight": "model-00119-of-000163.safetensors", + "model.layers.46.self_attn.q_a_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.self_attn.q_a_layernorm.weight": "model-00119-of-000163.safetensors", + "model.layers.46.self_attn.q_b_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.self_attn.kv_a_proj_with_mqa.weight": "model-00119-of-000163.safetensors", + "model.layers.46.self_attn.kv_a_layernorm.weight": "model-00119-of-000163.safetensors", + "model.layers.46.self_attn.kv_b_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.self_attn.o_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.gate.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.gate.e_score_correction_bias": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.shared_experts.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.shared_experts.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.shared_experts.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.0.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.0.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.0.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.1.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.1.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.1.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.2.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.2.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.2.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.3.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.3.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.3.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.4.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.4.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.4.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.5.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.5.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.5.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.6.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.6.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.6.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.7.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.7.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.7.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.8.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.8.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.8.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.9.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.9.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.9.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.10.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.10.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.10.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.11.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.11.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.11.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.12.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.12.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.12.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.13.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.13.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.13.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.14.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.14.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.14.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.15.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.15.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.15.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.16.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.16.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.16.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.17.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.17.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.17.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.18.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.18.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.18.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.19.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.19.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.19.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.20.gate_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.20.up_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.20.down_proj.weight": "model-00119-of-000163.safetensors", + "model.layers.46.mlp.experts.21.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.21.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.21.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.22.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.22.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.22.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.23.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.23.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.23.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.24.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.24.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.24.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.25.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.25.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.25.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.26.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.26.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.26.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.27.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.27.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.27.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.28.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.28.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.28.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.29.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.29.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.29.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.30.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.30.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.30.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.31.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.31.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.31.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.32.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.32.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.32.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.33.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.33.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.33.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.34.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.34.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.34.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.35.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.35.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.35.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.36.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.36.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.36.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.37.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.37.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.37.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.38.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.38.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.38.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.39.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.39.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.39.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.40.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.40.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.40.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.41.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.41.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.41.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.42.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.42.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.42.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.43.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.43.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.43.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.44.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.44.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.44.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.45.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.45.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.45.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.46.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.46.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.46.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.47.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.47.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.47.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.48.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.48.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.48.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.49.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.49.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.49.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.50.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.50.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.50.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.51.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.51.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.51.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.52.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.52.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.52.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.53.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.53.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.53.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.54.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.54.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.54.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.55.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.55.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.55.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.56.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.56.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.56.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.57.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.57.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.57.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.58.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.58.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.58.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.59.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.59.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.59.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.60.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.60.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.60.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.61.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.61.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.61.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.62.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.62.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.62.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.63.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.63.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.63.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.64.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.64.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.64.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.65.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.65.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.65.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.66.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.66.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.66.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.67.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.67.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.67.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.68.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.68.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.68.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.69.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.69.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.69.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.70.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.70.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.70.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.71.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.71.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.71.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.72.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.72.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.72.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.73.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.73.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.73.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.74.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.74.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.74.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.75.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.75.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.75.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.76.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.76.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.76.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.77.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.77.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.77.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.78.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.78.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.78.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.79.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.79.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.79.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.80.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.80.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.80.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.81.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.81.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.81.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.82.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.82.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.82.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.83.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.83.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.83.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.84.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.84.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.84.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.85.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.85.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.85.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.86.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.86.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.86.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.87.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.87.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.87.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.88.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.88.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.88.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.89.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.89.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.89.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.90.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.90.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.90.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.91.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.91.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.91.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.92.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.92.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.92.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.93.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.93.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.93.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.94.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.94.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.94.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.95.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.95.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.95.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.96.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.96.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.96.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.97.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.97.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.97.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.98.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.98.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.98.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.99.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.99.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.99.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.100.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.100.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.100.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.101.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.101.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.101.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.102.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.102.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.102.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.103.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.103.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.103.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.104.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.104.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.104.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.105.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.105.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.105.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.106.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.106.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.106.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.107.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.107.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.107.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.108.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.108.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.108.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.109.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.109.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.109.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.110.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.110.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.110.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.111.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.111.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.111.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.112.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.112.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.112.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.113.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.113.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.113.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.114.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.114.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.114.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.115.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.115.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.115.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.116.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.116.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.116.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.117.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.117.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.117.down_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.118.gate_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.118.up_proj.weight": "model-00120-of-000163.safetensors", + "model.layers.46.mlp.experts.118.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.119.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.119.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.119.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.120.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.120.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.120.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.121.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.121.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.121.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.122.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.122.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.122.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.123.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.123.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.123.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.124.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.124.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.124.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.125.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.125.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.125.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.126.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.126.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.126.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.127.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.127.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.127.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.128.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.128.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.128.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.129.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.129.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.129.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.130.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.130.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.130.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.131.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.131.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.131.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.132.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.132.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.132.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.133.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.133.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.133.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.134.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.134.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.134.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.135.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.135.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.135.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.136.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.136.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.136.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.137.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.137.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.137.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.138.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.138.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.138.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.139.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.139.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.139.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.140.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.140.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.140.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.141.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.141.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.141.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.142.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.142.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.142.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.143.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.143.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.143.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.144.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.144.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.144.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.145.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.145.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.145.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.146.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.146.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.146.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.147.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.147.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.147.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.148.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.148.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.148.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.149.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.149.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.149.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.150.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.150.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.150.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.151.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.151.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.151.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.152.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.152.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.152.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.153.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.153.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.153.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.154.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.154.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.154.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.155.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.155.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.155.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.156.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.156.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.156.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.157.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.157.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.157.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.158.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.158.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.158.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.159.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.159.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.159.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.160.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.160.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.160.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.161.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.161.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.161.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.162.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.162.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.162.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.163.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.163.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.163.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.164.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.164.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.164.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.165.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.165.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.165.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.166.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.166.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.166.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.167.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.167.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.167.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.168.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.168.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.168.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.169.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.169.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.169.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.170.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.170.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.170.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.171.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.171.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.171.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.172.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.172.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.172.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.173.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.173.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.173.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.174.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.174.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.174.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.175.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.175.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.175.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.176.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.176.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.176.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.177.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.177.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.177.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.178.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.178.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.178.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.179.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.179.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.179.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.180.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.180.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.180.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.181.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.181.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.181.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.182.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.182.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.182.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.183.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.183.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.183.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.184.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.184.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.184.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.185.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.185.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.185.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.186.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.186.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.186.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.187.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.187.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.187.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.188.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.188.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.188.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.189.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.189.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.189.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.190.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.190.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.190.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.191.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.191.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.191.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.192.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.192.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.192.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.193.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.193.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.193.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.194.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.194.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.194.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.195.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.195.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.195.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.196.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.196.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.196.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.197.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.197.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.197.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.198.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.198.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.198.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.199.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.199.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.199.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.200.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.200.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.200.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.201.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.201.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.201.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.202.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.202.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.202.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.203.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.203.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.203.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.204.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.204.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.204.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.205.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.205.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.205.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.206.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.206.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.206.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.207.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.207.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.207.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.208.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.208.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.208.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.209.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.209.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.209.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.210.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.210.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.210.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.211.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.211.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.211.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.212.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.212.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.212.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.213.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.213.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.213.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.214.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.214.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.214.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.215.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.215.up_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.215.down_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.216.gate_proj.weight": "model-00121-of-000163.safetensors", + "model.layers.46.mlp.experts.216.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.216.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.217.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.217.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.217.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.218.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.218.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.218.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.219.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.219.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.219.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.220.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.220.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.220.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.221.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.221.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.221.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.222.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.222.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.222.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.223.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.223.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.223.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.224.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.224.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.224.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.225.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.225.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.225.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.226.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.226.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.226.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.227.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.227.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.227.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.228.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.228.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.228.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.229.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.229.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.229.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.230.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.230.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.230.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.231.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.231.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.231.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.232.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.232.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.232.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.233.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.233.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.233.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.234.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.234.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.234.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.235.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.235.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.235.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.236.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.236.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.236.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.237.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.237.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.237.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.238.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.238.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.238.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.239.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.239.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.239.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.240.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.240.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.240.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.241.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.241.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.241.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.242.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.242.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.242.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.243.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.243.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.243.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.244.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.244.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.244.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.245.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.245.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.245.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.246.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.246.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.246.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.247.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.247.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.247.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.248.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.248.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.248.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.249.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.249.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.249.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.250.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.250.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.250.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.251.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.251.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.251.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.252.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.252.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.252.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.253.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.253.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.253.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.254.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.254.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.254.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.255.gate_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.255.up_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.mlp.experts.255.down_proj.weight": "model-00122-of-000163.safetensors", + "model.layers.46.input_layernorm.weight": "model-00122-of-000163.safetensors", + "model.layers.46.post_attention_layernorm.weight": "model-00122-of-000163.safetensors", + "model.layers.47.self_attn.q_a_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.self_attn.q_a_layernorm.weight": "model-00123-of-000163.safetensors", + "model.layers.47.self_attn.q_b_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.self_attn.kv_a_proj_with_mqa.weight": "model-00123-of-000163.safetensors", + "model.layers.47.self_attn.kv_a_layernorm.weight": "model-00123-of-000163.safetensors", + "model.layers.47.self_attn.kv_b_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.self_attn.o_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.gate.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.gate.e_score_correction_bias": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.shared_experts.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.shared_experts.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.shared_experts.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.0.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.0.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.0.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.1.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.1.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.1.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.2.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.2.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.2.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.3.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.3.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.3.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.4.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.4.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.4.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.5.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.5.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.5.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.6.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.6.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.6.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.7.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.7.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.7.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.8.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.8.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.8.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.9.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.9.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.9.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.10.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.10.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.10.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.11.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.11.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.11.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.12.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.12.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.12.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.13.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.13.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.13.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.14.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.14.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.14.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.15.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.15.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.15.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.16.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.16.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.16.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.17.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.17.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.17.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.18.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.18.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.18.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.19.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.19.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.19.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.20.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.20.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.20.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.21.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.21.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.21.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.22.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.22.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.22.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.23.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.23.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.23.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.24.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.24.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.24.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.25.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.25.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.25.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.26.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.26.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.26.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.27.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.27.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.27.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.28.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.28.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.28.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.29.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.29.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.29.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.30.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.30.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.30.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.31.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.31.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.31.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.32.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.32.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.32.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.33.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.33.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.33.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.34.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.34.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.34.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.35.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.35.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.35.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.36.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.36.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.36.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.37.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.37.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.37.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.38.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.38.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.38.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.39.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.39.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.39.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.40.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.40.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.40.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.41.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.41.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.41.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.42.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.42.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.42.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.43.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.43.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.43.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.44.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.44.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.44.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.45.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.45.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.45.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.46.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.46.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.46.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.47.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.47.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.47.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.48.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.48.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.48.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.49.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.49.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.49.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.50.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.50.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.50.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.51.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.51.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.51.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.52.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.52.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.52.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.53.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.53.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.53.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.54.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.54.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.54.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.55.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.55.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.55.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.56.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.56.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.56.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.57.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.57.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.57.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.58.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.58.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.58.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.59.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.59.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.59.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.60.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.60.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.60.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.61.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.61.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.61.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.62.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.62.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.62.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.63.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.63.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.63.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.64.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.64.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.64.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.65.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.65.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.65.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.66.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.66.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.66.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.67.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.67.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.67.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.68.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.68.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.68.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.69.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.69.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.69.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.70.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.70.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.70.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.71.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.71.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.71.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.72.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.72.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.72.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.73.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.73.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.73.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.74.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.74.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.74.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.75.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.75.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.75.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.76.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.76.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.76.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.77.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.77.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.77.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.78.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.78.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.78.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.79.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.79.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.79.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.80.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.80.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.80.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.81.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.81.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.81.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.82.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.82.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.82.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.83.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.83.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.83.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.84.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.84.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.84.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.85.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.85.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.85.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.86.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.86.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.86.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.87.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.87.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.87.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.88.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.88.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.88.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.89.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.89.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.89.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.90.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.90.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.90.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.91.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.91.up_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.91.down_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.92.gate_proj.weight": "model-00123-of-000163.safetensors", + "model.layers.47.mlp.experts.92.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.92.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.93.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.93.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.93.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.94.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.94.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.94.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.95.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.95.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.95.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.96.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.96.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.96.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.97.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.97.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.97.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.98.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.98.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.98.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.99.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.99.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.99.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.100.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.100.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.100.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.101.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.101.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.101.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.102.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.102.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.102.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.103.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.103.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.103.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.104.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.104.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.104.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.105.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.105.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.105.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.106.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.106.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.106.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.107.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.107.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.107.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.108.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.108.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.108.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.109.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.109.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.109.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.110.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.110.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.110.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.111.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.111.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.111.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.112.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.112.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.112.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.113.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.113.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.113.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.114.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.114.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.114.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.115.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.115.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.115.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.116.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.116.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.116.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.117.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.117.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.117.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.118.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.118.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.118.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.119.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.119.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.119.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.120.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.120.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.120.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.121.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.121.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.121.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.122.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.122.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.122.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.123.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.123.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.123.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.124.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.124.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.124.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.125.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.125.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.125.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.126.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.126.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.126.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.127.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.127.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.127.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.128.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.128.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.128.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.129.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.129.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.129.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.130.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.130.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.130.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.131.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.131.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.131.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.132.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.132.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.132.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.133.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.133.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.133.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.134.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.134.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.134.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.135.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.135.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.135.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.136.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.136.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.136.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.137.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.137.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.137.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.138.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.138.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.138.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.139.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.139.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.139.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.140.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.140.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.140.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.141.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.141.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.141.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.142.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.142.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.142.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.143.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.143.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.143.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.144.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.144.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.144.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.145.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.145.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.145.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.146.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.146.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.146.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.147.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.147.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.147.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.148.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.148.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.148.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.149.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.149.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.149.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.150.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.150.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.150.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.151.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.151.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.151.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.152.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.152.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.152.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.153.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.153.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.153.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.154.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.154.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.154.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.155.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.155.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.155.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.156.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.156.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.156.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.157.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.157.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.157.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.158.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.158.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.158.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.159.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.159.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.159.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.160.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.160.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.160.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.161.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.161.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.161.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.162.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.162.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.162.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.163.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.163.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.163.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.164.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.164.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.164.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.165.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.165.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.165.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.166.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.166.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.166.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.167.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.167.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.167.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.168.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.168.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.168.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.169.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.169.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.169.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.170.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.170.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.170.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.171.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.171.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.171.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.172.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.172.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.172.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.173.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.173.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.173.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.174.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.174.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.174.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.175.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.175.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.175.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.176.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.176.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.176.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.177.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.177.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.177.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.178.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.178.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.178.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.179.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.179.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.179.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.180.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.180.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.180.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.181.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.181.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.181.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.182.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.182.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.182.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.183.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.183.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.183.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.184.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.184.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.184.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.185.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.185.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.185.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.186.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.186.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.186.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.187.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.187.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.187.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.188.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.188.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.188.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.189.gate_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.189.up_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.189.down_proj.weight": "model-00124-of-000163.safetensors", + "model.layers.47.mlp.experts.190.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.190.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.190.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.191.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.191.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.191.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.192.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.192.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.192.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.193.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.193.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.193.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.194.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.194.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.194.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.195.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.195.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.195.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.196.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.196.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.196.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.197.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.197.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.197.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.198.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.198.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.198.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.199.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.199.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.199.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.200.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.200.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.200.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.201.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.201.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.201.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.202.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.202.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.202.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.203.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.203.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.203.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.204.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.204.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.204.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.205.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.205.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.205.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.206.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.206.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.206.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.207.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.207.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.207.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.208.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.208.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.208.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.209.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.209.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.209.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.210.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.210.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.210.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.211.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.211.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.211.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.212.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.212.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.212.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.213.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.213.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.213.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.214.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.214.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.214.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.215.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.215.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.215.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.216.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.216.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.216.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.217.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.217.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.217.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.218.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.218.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.218.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.219.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.219.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.219.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.220.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.220.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.220.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.221.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.221.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.221.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.222.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.222.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.222.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.223.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.223.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.223.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.224.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.224.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.224.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.225.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.225.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.225.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.226.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.226.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.226.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.227.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.227.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.227.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.228.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.228.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.228.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.229.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.229.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.229.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.230.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.230.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.230.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.231.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.231.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.231.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.232.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.232.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.232.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.233.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.233.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.233.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.234.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.234.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.234.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.235.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.235.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.235.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.236.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.236.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.236.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.237.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.237.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.237.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.238.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.238.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.238.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.239.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.239.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.239.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.240.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.240.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.240.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.241.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.241.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.241.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.242.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.242.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.242.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.243.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.243.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.243.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.244.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.244.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.244.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.245.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.245.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.245.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.246.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.246.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.246.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.247.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.247.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.247.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.248.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.248.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.248.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.249.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.249.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.249.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.250.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.250.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.250.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.251.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.251.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.251.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.252.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.252.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.252.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.253.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.253.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.253.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.254.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.254.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.254.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.255.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.255.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.mlp.experts.255.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.47.input_layernorm.weight": "model-00125-of-000163.safetensors", + "model.layers.47.post_attention_layernorm.weight": "model-00125-of-000163.safetensors", + "model.layers.48.self_attn.q_a_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.self_attn.q_a_layernorm.weight": "model-00125-of-000163.safetensors", + "model.layers.48.self_attn.q_b_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.self_attn.kv_a_proj_with_mqa.weight": "model-00125-of-000163.safetensors", + "model.layers.48.self_attn.kv_a_layernorm.weight": "model-00125-of-000163.safetensors", + "model.layers.48.self_attn.kv_b_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.self_attn.o_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.gate.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.gate.e_score_correction_bias": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.shared_experts.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.shared_experts.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.shared_experts.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.0.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.0.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.0.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.1.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.1.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.1.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.2.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.2.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.2.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.3.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.3.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.3.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.4.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.4.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.4.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.5.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.5.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.5.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.6.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.6.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.6.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.7.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.7.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.7.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.8.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.8.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.8.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.9.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.9.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.9.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.10.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.10.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.10.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.11.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.11.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.11.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.12.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.12.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.12.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.13.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.13.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.13.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.14.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.14.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.14.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.15.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.15.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.15.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.16.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.16.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.16.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.17.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.17.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.17.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.18.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.18.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.18.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.19.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.19.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.19.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.20.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.20.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.20.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.21.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.21.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.21.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.22.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.22.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.22.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.23.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.23.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.23.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.24.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.24.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.24.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.25.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.25.up_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.25.down_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.26.gate_proj.weight": "model-00125-of-000163.safetensors", + "model.layers.48.mlp.experts.26.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.26.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.27.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.27.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.27.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.28.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.28.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.28.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.29.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.29.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.29.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.30.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.30.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.30.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.31.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.31.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.31.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.32.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.32.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.32.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.33.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.33.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.33.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.34.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.34.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.34.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.35.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.35.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.35.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.36.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.36.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.36.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.37.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.37.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.37.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.38.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.38.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.38.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.39.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.39.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.39.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.40.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.40.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.40.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.41.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.41.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.41.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.42.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.42.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.42.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.43.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.43.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.43.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.44.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.44.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.44.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.45.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.45.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.45.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.46.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.46.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.46.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.47.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.47.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.47.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.48.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.48.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.48.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.49.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.49.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.49.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.50.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.50.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.50.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.51.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.51.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.51.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.52.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.52.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.52.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.53.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.53.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.53.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.54.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.54.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.54.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.55.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.55.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.55.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.56.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.56.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.56.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.57.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.57.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.57.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.58.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.58.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.58.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.59.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.59.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.59.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.60.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.60.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.60.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.61.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.61.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.61.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.62.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.62.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.62.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.63.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.63.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.63.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.64.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.64.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.64.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.65.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.65.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.65.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.66.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.66.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.66.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.67.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.67.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.67.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.68.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.68.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.68.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.69.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.69.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.69.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.70.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.70.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.70.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.71.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.71.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.71.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.72.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.72.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.72.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.73.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.73.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.73.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.74.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.74.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.74.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.75.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.75.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.75.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.76.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.76.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.76.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.77.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.77.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.77.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.78.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.78.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.78.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.79.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.79.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.79.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.80.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.80.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.80.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.81.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.81.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.81.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.82.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.82.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.82.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.83.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.83.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.83.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.84.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.84.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.84.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.85.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.85.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.85.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.86.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.86.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.86.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.87.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.87.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.87.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.88.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.88.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.88.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.89.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.89.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.89.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.90.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.90.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.90.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.91.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.91.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.91.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.92.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.92.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.92.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.93.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.93.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.93.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.94.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.94.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.94.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.95.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.95.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.95.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.96.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.96.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.96.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.97.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.97.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.97.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.98.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.98.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.98.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.99.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.99.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.99.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.100.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.100.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.100.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.101.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.101.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.101.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.102.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.102.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.102.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.103.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.103.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.103.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.104.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.104.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.104.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.105.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.105.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.105.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.106.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.106.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.106.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.107.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.107.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.107.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.108.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.108.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.108.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.109.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.109.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.109.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.110.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.110.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.110.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.111.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.111.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.111.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.112.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.112.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.112.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.113.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.113.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.113.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.114.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.114.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.114.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.115.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.115.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.115.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.116.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.116.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.116.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.117.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.117.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.117.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.118.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.118.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.118.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.119.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.119.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.119.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.120.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.120.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.120.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.121.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.121.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.121.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.122.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.122.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.122.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.123.gate_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.123.up_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.123.down_proj.weight": "model-00126-of-000163.safetensors", + "model.layers.48.mlp.experts.124.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.124.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.124.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.125.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.125.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.125.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.126.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.126.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.126.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.127.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.127.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.127.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.128.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.128.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.128.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.129.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.129.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.129.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.130.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.130.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.130.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.131.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.131.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.131.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.132.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.132.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.132.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.133.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.133.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.133.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.134.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.134.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.134.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.135.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.135.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.135.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.136.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.136.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.136.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.137.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.137.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.137.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.138.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.138.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.138.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.139.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.139.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.139.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.140.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.140.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.140.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.141.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.141.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.141.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.142.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.142.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.142.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.143.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.143.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.143.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.144.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.144.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.144.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.145.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.145.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.145.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.146.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.146.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.146.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.147.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.147.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.147.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.148.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.148.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.148.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.149.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.149.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.149.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.150.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.150.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.150.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.151.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.151.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.151.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.152.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.152.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.152.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.153.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.153.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.153.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.154.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.154.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.154.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.155.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.155.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.155.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.156.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.156.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.156.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.157.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.157.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.157.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.158.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.158.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.158.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.159.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.159.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.159.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.160.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.160.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.160.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.161.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.161.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.161.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.162.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.162.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.162.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.163.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.163.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.163.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.164.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.164.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.164.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.165.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.165.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.165.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.166.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.166.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.166.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.167.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.167.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.167.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.168.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.168.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.168.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.169.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.169.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.169.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.170.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.170.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.170.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.171.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.171.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.171.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.172.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.172.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.172.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.173.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.173.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.173.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.174.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.174.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.174.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.175.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.175.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.175.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.176.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.176.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.176.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.177.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.177.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.177.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.178.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.178.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.178.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.179.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.179.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.179.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.180.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.180.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.180.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.181.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.181.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.181.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.182.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.182.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.182.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.183.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.183.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.183.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.184.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.184.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.184.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.185.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.185.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.185.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.186.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.186.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.186.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.187.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.187.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.187.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.188.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.188.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.188.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.189.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.189.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.189.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.190.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.190.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.190.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.191.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.191.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.191.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.192.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.192.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.192.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.193.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.193.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.193.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.194.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.194.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.194.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.195.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.195.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.195.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.196.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.196.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.196.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.197.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.197.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.197.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.198.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.198.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.198.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.199.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.199.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.199.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.200.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.200.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.200.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.201.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.201.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.201.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.202.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.202.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.202.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.203.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.203.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.203.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.204.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.204.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.204.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.205.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.205.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.205.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.206.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.206.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.206.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.207.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.207.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.207.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.208.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.208.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.208.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.209.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.209.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.209.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.210.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.210.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.210.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.211.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.211.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.211.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.212.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.212.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.212.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.213.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.213.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.213.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.214.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.214.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.214.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.215.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.215.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.215.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.216.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.216.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.216.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.217.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.217.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.217.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.218.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.218.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.218.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.219.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.219.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.219.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.220.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.220.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.220.down_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.221.gate_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.221.up_proj.weight": "model-00127-of-000163.safetensors", + "model.layers.48.mlp.experts.221.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.222.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.222.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.222.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.223.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.223.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.223.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.224.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.224.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.224.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.225.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.225.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.225.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.226.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.226.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.226.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.227.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.227.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.227.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.228.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.228.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.228.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.229.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.229.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.229.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.230.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.230.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.230.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.231.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.231.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.231.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.232.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.232.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.232.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.233.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.233.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.233.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.234.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.234.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.234.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.235.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.235.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.235.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.236.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.236.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.236.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.237.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.237.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.237.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.238.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.238.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.238.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.239.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.239.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.239.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.240.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.240.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.240.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.241.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.241.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.241.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.242.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.242.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.242.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.243.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.243.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.243.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.244.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.244.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.244.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.245.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.245.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.245.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.246.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.246.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.246.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.247.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.247.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.247.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.248.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.248.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.248.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.249.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.249.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.249.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.250.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.250.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.250.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.251.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.251.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.251.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.252.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.252.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.252.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.253.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.253.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.253.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.254.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.254.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.254.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.255.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.255.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.mlp.experts.255.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.48.input_layernorm.weight": "model-00128-of-000163.safetensors", + "model.layers.48.post_attention_layernorm.weight": "model-00128-of-000163.safetensors", + "model.layers.49.self_attn.q_a_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.self_attn.q_a_layernorm.weight": "model-00128-of-000163.safetensors", + "model.layers.49.self_attn.q_b_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.self_attn.kv_a_proj_with_mqa.weight": "model-00128-of-000163.safetensors", + "model.layers.49.self_attn.kv_a_layernorm.weight": "model-00128-of-000163.safetensors", + "model.layers.49.self_attn.kv_b_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.self_attn.o_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.gate.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.gate.e_score_correction_bias": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.shared_experts.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.shared_experts.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.shared_experts.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.0.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.0.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.0.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.1.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.1.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.1.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.2.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.2.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.2.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.3.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.3.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.3.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.4.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.4.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.4.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.5.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.5.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.5.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.6.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.6.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.6.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.7.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.7.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.7.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.8.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.8.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.8.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.9.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.9.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.9.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.10.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.10.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.10.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.11.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.11.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.11.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.12.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.12.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.12.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.13.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.13.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.13.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.14.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.14.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.14.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.15.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.15.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.15.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.16.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.16.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.16.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.17.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.17.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.17.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.18.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.18.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.18.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.19.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.19.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.19.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.20.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.20.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.20.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.21.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.21.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.21.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.22.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.22.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.22.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.23.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.23.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.23.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.24.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.24.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.24.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.25.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.25.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.25.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.26.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.26.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.26.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.27.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.27.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.27.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.28.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.28.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.28.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.29.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.29.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.29.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.30.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.30.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.30.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.31.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.31.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.31.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.32.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.32.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.32.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.33.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.33.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.33.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.34.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.34.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.34.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.35.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.35.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.35.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.36.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.36.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.36.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.37.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.37.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.37.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.38.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.38.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.38.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.39.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.39.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.39.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.40.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.40.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.40.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.41.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.41.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.41.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.42.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.42.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.42.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.43.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.43.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.43.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.44.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.44.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.44.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.45.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.45.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.45.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.46.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.46.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.46.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.47.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.47.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.47.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.48.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.48.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.48.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.49.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.49.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.49.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.50.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.50.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.50.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.51.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.51.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.51.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.52.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.52.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.52.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.53.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.53.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.53.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.54.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.54.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.54.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.55.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.55.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.55.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.56.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.56.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.56.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.57.gate_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.57.up_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.57.down_proj.weight": "model-00128-of-000163.safetensors", + "model.layers.49.mlp.experts.58.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.58.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.58.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.59.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.59.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.59.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.60.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.60.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.60.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.61.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.61.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.61.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.62.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.62.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.62.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.63.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.63.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.63.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.64.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.64.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.64.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.65.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.65.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.65.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.66.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.66.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.66.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.67.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.67.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.67.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.68.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.68.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.68.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.69.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.69.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.69.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.70.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.70.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.70.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.71.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.71.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.71.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.72.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.72.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.72.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.73.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.73.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.73.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.74.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.74.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.74.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.75.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.75.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.75.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.76.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.76.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.76.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.77.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.77.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.77.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.78.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.78.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.78.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.79.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.79.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.79.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.80.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.80.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.80.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.81.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.81.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.81.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.82.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.82.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.82.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.83.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.83.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.83.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.84.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.84.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.84.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.85.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.85.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.85.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.86.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.86.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.86.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.87.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.87.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.87.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.88.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.88.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.88.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.89.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.89.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.89.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.90.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.90.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.90.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.91.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.91.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.91.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.92.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.92.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.92.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.93.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.93.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.93.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.94.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.94.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.94.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.95.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.95.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.95.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.96.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.96.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.96.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.97.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.97.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.97.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.98.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.98.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.98.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.99.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.99.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.99.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.100.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.100.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.100.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.101.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.101.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.101.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.102.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.102.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.102.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.103.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.103.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.103.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.104.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.104.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.104.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.105.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.105.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.105.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.106.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.106.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.106.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.107.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.107.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.107.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.108.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.108.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.108.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.109.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.109.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.109.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.110.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.110.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.110.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.111.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.111.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.111.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.112.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.112.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.112.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.113.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.113.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.113.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.114.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.114.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.114.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.115.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.115.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.115.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.116.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.116.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.116.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.117.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.117.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.117.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.118.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.118.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.118.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.119.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.119.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.119.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.120.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.120.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.120.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.121.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.121.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.121.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.122.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.122.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.122.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.123.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.123.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.123.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.124.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.124.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.124.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.125.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.125.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.125.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.126.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.126.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.126.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.127.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.127.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.127.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.128.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.128.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.128.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.129.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.129.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.129.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.130.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.130.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.130.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.131.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.131.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.131.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.132.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.132.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.132.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.133.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.133.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.133.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.134.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.134.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.134.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.135.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.135.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.135.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.136.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.136.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.136.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.137.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.137.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.137.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.138.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.138.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.138.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.139.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.139.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.139.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.140.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.140.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.140.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.141.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.141.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.141.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.142.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.142.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.142.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.143.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.143.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.143.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.144.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.144.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.144.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.145.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.145.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.145.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.146.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.146.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.146.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.147.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.147.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.147.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.148.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.148.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.148.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.149.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.149.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.149.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.150.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.150.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.150.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.151.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.151.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.151.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.152.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.152.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.152.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.153.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.153.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.153.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.154.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.154.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.154.down_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.155.gate_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.155.up_proj.weight": "model-00129-of-000163.safetensors", + "model.layers.49.mlp.experts.155.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.156.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.156.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.156.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.157.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.157.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.157.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.158.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.158.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.158.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.159.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.159.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.159.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.160.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.160.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.160.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.161.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.161.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.161.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.162.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.162.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.162.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.163.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.163.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.163.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.164.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.164.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.164.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.165.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.165.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.165.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.166.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.166.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.166.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.167.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.167.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.167.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.168.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.168.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.168.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.169.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.169.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.169.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.170.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.170.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.170.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.171.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.171.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.171.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.172.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.172.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.172.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.173.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.173.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.173.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.174.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.174.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.174.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.175.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.175.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.175.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.176.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.176.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.176.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.177.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.177.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.177.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.178.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.178.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.178.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.179.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.179.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.179.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.180.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.180.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.180.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.181.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.181.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.181.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.182.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.182.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.182.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.183.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.183.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.183.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.184.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.184.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.184.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.185.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.185.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.185.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.186.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.186.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.186.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.187.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.187.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.187.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.188.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.188.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.188.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.189.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.189.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.189.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.190.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.190.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.190.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.191.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.191.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.191.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.192.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.192.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.192.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.193.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.193.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.193.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.194.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.194.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.194.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.195.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.195.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.195.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.196.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.196.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.196.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.197.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.197.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.197.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.198.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.198.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.198.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.199.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.199.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.199.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.200.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.200.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.200.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.201.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.201.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.201.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.202.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.202.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.202.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.203.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.203.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.203.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.204.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.204.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.204.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.205.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.205.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.205.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.206.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.206.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.206.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.207.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.207.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.207.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.208.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.208.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.208.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.209.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.209.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.209.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.210.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.210.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.210.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.211.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.211.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.211.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.212.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.212.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.212.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.213.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.213.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.213.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.214.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.214.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.214.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.215.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.215.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.215.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.216.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.216.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.216.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.217.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.217.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.217.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.218.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.218.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.218.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.219.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.219.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.219.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.220.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.220.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.220.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.221.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.221.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.221.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.222.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.222.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.222.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.223.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.223.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.223.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.224.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.224.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.224.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.225.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.225.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.225.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.226.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.226.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.226.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.227.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.227.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.227.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.228.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.228.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.228.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.229.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.229.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.229.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.230.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.230.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.230.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.231.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.231.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.231.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.232.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.232.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.232.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.233.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.233.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.233.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.234.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.234.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.234.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.235.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.235.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.235.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.236.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.236.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.236.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.237.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.237.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.237.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.238.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.238.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.238.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.239.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.239.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.239.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.240.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.240.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.240.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.241.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.241.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.241.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.242.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.242.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.242.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.243.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.243.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.243.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.244.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.244.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.244.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.245.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.245.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.245.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.246.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.246.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.246.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.247.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.247.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.247.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.248.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.248.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.248.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.249.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.249.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.249.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.250.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.250.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.250.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.251.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.251.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.251.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.252.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.252.up_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.252.down_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.253.gate_proj.weight": "model-00130-of-000163.safetensors", + "model.layers.49.mlp.experts.253.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.49.mlp.experts.253.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.49.mlp.experts.254.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.49.mlp.experts.254.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.49.mlp.experts.254.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.49.mlp.experts.255.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.49.mlp.experts.255.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.49.mlp.experts.255.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.49.input_layernorm.weight": "model-00131-of-000163.safetensors", + "model.layers.49.post_attention_layernorm.weight": "model-00131-of-000163.safetensors", + "model.layers.50.self_attn.q_a_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.self_attn.q_a_layernorm.weight": "model-00131-of-000163.safetensors", + "model.layers.50.self_attn.q_b_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.self_attn.kv_a_proj_with_mqa.weight": "model-00131-of-000163.safetensors", + "model.layers.50.self_attn.kv_a_layernorm.weight": "model-00131-of-000163.safetensors", + "model.layers.50.self_attn.kv_b_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.self_attn.o_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.gate.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.gate.e_score_correction_bias": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.shared_experts.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.shared_experts.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.shared_experts.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.0.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.0.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.0.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.1.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.1.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.1.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.2.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.2.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.2.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.3.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.3.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.3.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.4.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.4.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.4.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.5.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.5.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.5.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.6.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.6.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.6.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.7.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.7.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.7.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.8.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.8.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.8.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.9.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.9.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.9.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.10.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.10.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.10.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.11.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.11.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.11.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.12.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.12.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.12.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.13.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.13.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.13.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.14.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.14.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.14.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.15.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.15.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.15.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.16.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.16.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.16.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.17.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.17.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.17.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.18.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.18.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.18.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.19.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.19.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.19.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.20.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.20.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.20.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.21.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.21.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.21.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.22.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.22.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.22.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.23.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.23.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.23.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.24.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.24.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.24.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.25.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.25.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.25.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.26.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.26.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.26.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.27.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.27.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.27.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.28.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.28.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.28.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.29.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.29.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.29.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.30.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.30.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.30.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.31.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.31.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.31.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.32.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.32.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.32.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.33.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.33.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.33.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.34.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.34.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.34.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.35.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.35.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.35.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.36.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.36.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.36.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.37.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.37.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.37.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.38.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.38.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.38.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.39.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.39.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.39.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.40.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.40.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.40.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.41.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.41.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.41.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.42.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.42.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.42.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.43.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.43.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.43.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.44.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.44.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.44.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.45.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.45.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.45.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.46.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.46.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.46.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.47.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.47.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.47.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.48.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.48.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.48.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.49.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.49.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.49.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.50.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.50.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.50.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.51.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.51.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.51.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.52.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.52.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.52.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.53.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.53.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.53.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.54.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.54.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.54.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.55.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.55.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.55.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.56.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.56.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.56.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.57.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.57.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.57.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.58.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.58.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.58.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.59.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.59.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.59.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.60.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.60.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.60.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.61.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.61.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.61.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.62.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.62.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.62.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.63.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.63.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.63.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.64.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.64.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.64.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.65.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.65.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.65.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.66.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.66.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.66.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.67.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.67.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.67.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.68.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.68.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.68.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.69.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.69.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.69.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.70.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.70.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.70.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.71.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.71.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.71.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.72.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.72.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.72.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.73.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.73.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.73.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.74.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.74.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.74.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.75.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.75.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.75.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.76.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.76.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.76.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.77.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.77.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.77.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.78.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.78.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.78.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.79.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.79.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.79.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.80.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.80.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.80.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.81.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.81.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.81.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.82.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.82.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.82.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.83.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.83.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.83.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.84.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.84.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.84.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.85.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.85.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.85.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.86.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.86.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.86.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.87.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.87.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.87.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.88.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.88.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.88.down_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.89.gate_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.89.up_proj.weight": "model-00131-of-000163.safetensors", + "model.layers.50.mlp.experts.89.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.90.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.90.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.90.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.91.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.91.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.91.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.92.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.92.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.92.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.93.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.93.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.93.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.94.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.94.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.94.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.95.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.95.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.95.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.96.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.96.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.96.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.97.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.97.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.97.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.98.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.98.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.98.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.99.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.99.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.99.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.100.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.100.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.100.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.101.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.101.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.101.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.102.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.102.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.102.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.103.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.103.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.103.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.104.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.104.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.104.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.105.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.105.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.105.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.106.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.106.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.106.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.107.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.107.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.107.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.108.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.108.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.108.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.109.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.109.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.109.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.110.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.110.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.110.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.111.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.111.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.111.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.112.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.112.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.112.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.113.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.113.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.113.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.114.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.114.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.114.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.115.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.115.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.115.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.116.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.116.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.116.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.117.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.117.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.117.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.118.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.118.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.118.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.119.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.119.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.119.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.120.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.120.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.120.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.121.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.121.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.121.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.122.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.122.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.122.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.123.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.123.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.123.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.124.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.124.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.124.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.125.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.125.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.125.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.126.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.126.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.126.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.127.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.127.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.127.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.128.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.128.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.128.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.129.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.129.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.129.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.130.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.130.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.130.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.131.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.131.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.131.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.132.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.132.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.132.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.133.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.133.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.133.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.134.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.134.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.134.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.135.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.135.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.135.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.136.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.136.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.136.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.137.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.137.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.137.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.138.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.138.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.138.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.139.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.139.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.139.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.140.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.140.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.140.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.141.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.141.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.141.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.142.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.142.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.142.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.143.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.143.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.143.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.144.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.144.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.144.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.145.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.145.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.145.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.146.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.146.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.146.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.147.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.147.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.147.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.148.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.148.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.148.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.149.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.149.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.149.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.150.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.150.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.150.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.151.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.151.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.151.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.152.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.152.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.152.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.153.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.153.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.153.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.154.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.154.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.154.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.155.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.155.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.155.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.156.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.156.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.156.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.157.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.157.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.157.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.158.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.158.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.158.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.159.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.159.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.159.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.160.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.160.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.160.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.161.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.161.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.161.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.162.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.162.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.162.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.163.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.163.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.163.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.164.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.164.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.164.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.165.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.165.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.165.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.166.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.166.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.166.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.167.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.167.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.167.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.168.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.168.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.168.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.169.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.169.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.169.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.170.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.170.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.170.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.171.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.171.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.171.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.172.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.172.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.172.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.173.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.173.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.173.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.174.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.174.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.174.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.175.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.175.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.175.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.176.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.176.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.176.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.177.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.177.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.177.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.178.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.178.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.178.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.179.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.179.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.179.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.180.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.180.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.180.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.181.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.181.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.181.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.182.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.182.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.182.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.183.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.183.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.183.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.184.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.184.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.184.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.185.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.185.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.185.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.186.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.186.up_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.186.down_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.187.gate_proj.weight": "model-00132-of-000163.safetensors", + "model.layers.50.mlp.experts.187.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.187.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.188.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.188.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.188.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.189.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.189.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.189.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.190.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.190.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.190.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.191.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.191.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.191.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.192.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.192.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.192.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.193.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.193.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.193.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.194.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.194.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.194.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.195.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.195.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.195.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.196.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.196.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.196.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.197.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.197.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.197.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.198.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.198.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.198.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.199.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.199.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.199.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.200.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.200.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.200.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.201.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.201.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.201.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.202.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.202.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.202.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.203.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.203.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.203.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.204.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.204.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.204.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.205.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.205.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.205.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.206.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.206.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.206.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.207.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.207.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.207.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.208.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.208.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.208.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.209.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.209.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.209.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.210.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.210.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.210.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.211.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.211.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.211.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.212.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.212.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.212.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.213.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.213.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.213.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.214.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.214.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.214.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.215.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.215.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.215.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.216.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.216.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.216.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.217.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.217.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.217.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.218.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.218.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.218.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.219.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.219.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.219.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.220.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.220.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.220.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.221.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.221.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.221.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.222.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.222.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.222.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.223.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.223.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.223.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.224.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.224.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.224.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.225.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.225.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.225.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.226.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.226.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.226.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.227.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.227.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.227.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.228.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.228.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.228.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.229.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.229.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.229.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.230.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.230.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.230.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.231.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.231.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.231.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.232.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.232.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.232.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.233.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.233.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.233.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.234.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.234.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.234.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.235.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.235.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.235.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.236.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.236.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.236.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.237.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.237.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.237.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.238.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.238.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.238.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.239.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.239.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.239.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.240.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.240.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.240.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.241.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.241.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.241.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.242.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.242.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.242.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.243.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.243.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.243.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.244.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.244.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.244.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.245.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.245.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.245.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.246.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.246.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.246.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.247.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.247.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.247.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.248.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.248.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.248.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.249.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.249.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.249.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.250.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.250.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.250.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.251.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.251.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.251.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.252.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.252.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.252.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.253.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.253.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.253.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.254.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.254.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.254.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.255.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.255.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.mlp.experts.255.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.50.input_layernorm.weight": "model-00133-of-000163.safetensors", + "model.layers.50.post_attention_layernorm.weight": "model-00133-of-000163.safetensors", + "model.layers.51.self_attn.q_a_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.self_attn.q_a_layernorm.weight": "model-00133-of-000163.safetensors", + "model.layers.51.self_attn.q_b_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.self_attn.kv_a_proj_with_mqa.weight": "model-00133-of-000163.safetensors", + "model.layers.51.self_attn.kv_a_layernorm.weight": "model-00133-of-000163.safetensors", + "model.layers.51.self_attn.kv_b_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.self_attn.o_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.gate.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.gate.e_score_correction_bias": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.shared_experts.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.shared_experts.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.shared_experts.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.0.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.0.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.0.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.1.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.1.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.1.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.2.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.2.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.2.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.3.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.3.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.3.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.4.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.4.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.4.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.5.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.5.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.5.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.6.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.6.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.6.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.7.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.7.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.7.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.8.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.8.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.8.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.9.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.9.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.9.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.10.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.10.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.10.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.11.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.11.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.11.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.12.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.12.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.12.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.13.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.13.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.13.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.14.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.14.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.14.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.15.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.15.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.15.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.16.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.16.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.16.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.17.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.17.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.17.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.18.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.18.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.18.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.19.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.19.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.19.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.20.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.20.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.20.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.21.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.21.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.21.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.22.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.22.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.22.down_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.23.gate_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.23.up_proj.weight": "model-00133-of-000163.safetensors", + "model.layers.51.mlp.experts.23.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.24.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.24.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.24.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.25.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.25.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.25.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.26.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.26.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.26.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.27.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.27.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.27.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.28.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.28.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.28.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.29.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.29.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.29.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.30.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.30.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.30.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.31.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.31.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.31.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.32.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.32.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.32.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.33.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.33.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.33.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.34.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.34.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.34.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.35.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.35.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.35.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.36.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.36.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.36.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.37.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.37.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.37.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.38.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.38.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.38.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.39.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.39.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.39.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.40.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.40.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.40.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.41.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.41.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.41.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.42.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.42.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.42.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.43.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.43.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.43.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.44.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.44.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.44.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.45.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.45.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.45.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.46.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.46.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.46.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.47.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.47.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.47.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.48.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.48.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.48.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.49.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.49.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.49.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.50.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.50.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.50.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.51.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.51.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.51.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.52.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.52.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.52.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.53.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.53.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.53.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.54.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.54.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.54.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.55.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.55.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.55.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.56.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.56.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.56.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.57.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.57.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.57.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.58.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.58.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.58.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.59.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.59.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.59.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.60.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.60.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.60.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.61.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.61.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.61.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.62.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.62.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.62.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.63.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.63.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.63.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.64.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.64.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.64.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.65.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.65.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.65.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.66.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.66.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.66.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.67.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.67.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.67.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.68.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.68.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.68.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.69.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.69.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.69.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.70.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.70.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.70.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.71.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.71.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.71.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.72.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.72.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.72.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.73.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.73.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.73.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.74.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.74.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.74.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.75.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.75.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.75.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.76.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.76.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.76.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.77.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.77.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.77.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.78.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.78.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.78.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.79.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.79.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.79.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.80.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.80.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.80.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.81.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.81.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.81.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.82.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.82.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.82.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.83.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.83.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.83.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.84.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.84.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.84.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.85.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.85.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.85.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.86.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.86.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.86.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.87.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.87.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.87.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.88.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.88.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.88.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.89.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.89.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.89.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.90.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.90.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.90.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.91.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.91.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.91.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.92.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.92.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.92.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.93.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.93.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.93.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.94.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.94.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.94.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.95.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.95.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.95.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.96.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.96.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.96.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.97.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.97.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.97.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.98.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.98.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.98.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.99.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.99.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.99.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.100.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.100.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.100.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.101.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.101.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.101.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.102.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.102.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.102.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.103.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.103.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.103.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.104.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.104.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.104.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.105.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.105.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.105.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.106.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.106.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.106.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.107.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.107.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.107.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.108.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.108.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.108.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.109.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.109.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.109.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.110.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.110.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.110.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.111.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.111.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.111.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.112.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.112.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.112.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.113.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.113.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.113.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.114.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.114.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.114.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.115.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.115.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.115.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.116.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.116.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.116.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.117.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.117.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.117.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.118.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.118.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.118.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.119.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.119.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.119.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.120.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.120.up_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.120.down_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.121.gate_proj.weight": "model-00134-of-000163.safetensors", + "model.layers.51.mlp.experts.121.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.121.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.122.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.122.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.122.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.123.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.123.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.123.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.124.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.124.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.124.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.125.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.125.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.125.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.126.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.126.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.126.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.127.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.127.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.127.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.128.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.128.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.128.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.129.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.129.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.129.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.130.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.130.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.130.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.131.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.131.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.131.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.132.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.132.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.132.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.133.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.133.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.133.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.134.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.134.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.134.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.135.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.135.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.135.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.136.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.136.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.136.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.137.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.137.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.137.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.138.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.138.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.138.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.139.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.139.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.139.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.140.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.140.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.140.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.141.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.141.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.141.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.142.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.142.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.142.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.143.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.143.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.143.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.144.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.144.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.144.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.145.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.145.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.145.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.146.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.146.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.146.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.147.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.147.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.147.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.148.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.148.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.148.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.149.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.149.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.149.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.150.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.150.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.150.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.151.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.151.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.151.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.152.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.152.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.152.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.153.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.153.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.153.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.154.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.154.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.154.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.155.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.155.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.155.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.156.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.156.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.156.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.157.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.157.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.157.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.158.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.158.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.158.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.159.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.159.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.159.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.160.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.160.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.160.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.161.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.161.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.161.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.162.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.162.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.162.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.163.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.163.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.163.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.164.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.164.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.164.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.165.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.165.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.165.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.166.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.166.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.166.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.167.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.167.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.167.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.168.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.168.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.168.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.169.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.169.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.169.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.170.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.170.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.170.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.171.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.171.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.171.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.172.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.172.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.172.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.173.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.173.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.173.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.174.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.174.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.174.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.175.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.175.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.175.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.176.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.176.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.176.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.177.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.177.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.177.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.178.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.178.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.178.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.179.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.179.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.179.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.180.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.180.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.180.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.181.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.181.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.181.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.182.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.182.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.182.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.183.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.183.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.183.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.184.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.184.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.184.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.185.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.185.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.185.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.186.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.186.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.186.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.187.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.187.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.187.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.188.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.188.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.188.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.189.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.189.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.189.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.190.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.190.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.190.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.191.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.191.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.191.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.192.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.192.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.192.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.193.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.193.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.193.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.194.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.194.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.194.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.195.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.195.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.195.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.196.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.196.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.196.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.197.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.197.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.197.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.198.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.198.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.198.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.199.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.199.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.199.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.200.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.200.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.200.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.201.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.201.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.201.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.202.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.202.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.202.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.203.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.203.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.203.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.204.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.204.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.204.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.205.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.205.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.205.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.206.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.206.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.206.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.207.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.207.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.207.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.208.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.208.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.208.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.209.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.209.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.209.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.210.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.210.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.210.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.211.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.211.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.211.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.212.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.212.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.212.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.213.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.213.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.213.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.214.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.214.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.214.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.215.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.215.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.215.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.216.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.216.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.216.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.217.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.217.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.217.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.218.gate_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.218.up_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.218.down_proj.weight": "model-00135-of-000163.safetensors", + "model.layers.51.mlp.experts.219.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.219.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.219.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.220.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.220.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.220.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.221.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.221.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.221.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.222.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.222.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.222.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.223.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.223.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.223.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.224.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.224.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.224.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.225.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.225.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.225.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.226.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.226.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.226.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.227.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.227.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.227.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.228.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.228.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.228.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.229.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.229.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.229.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.230.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.230.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.230.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.231.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.231.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.231.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.232.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.232.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.232.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.233.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.233.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.233.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.234.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.234.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.234.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.235.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.235.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.235.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.236.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.236.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.236.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.237.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.237.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.237.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.238.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.238.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.238.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.239.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.239.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.239.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.240.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.240.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.240.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.241.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.241.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.241.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.242.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.242.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.242.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.243.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.243.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.243.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.244.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.244.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.244.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.245.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.245.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.245.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.246.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.246.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.246.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.247.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.247.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.247.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.248.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.248.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.248.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.249.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.249.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.249.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.250.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.250.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.250.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.251.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.251.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.251.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.252.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.252.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.252.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.253.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.253.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.253.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.254.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.254.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.254.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.255.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.255.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.mlp.experts.255.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.51.input_layernorm.weight": "model-00136-of-000163.safetensors", + "model.layers.51.post_attention_layernorm.weight": "model-00136-of-000163.safetensors", + "model.layers.52.self_attn.q_a_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.self_attn.q_a_layernorm.weight": "model-00136-of-000163.safetensors", + "model.layers.52.self_attn.q_b_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.self_attn.kv_a_proj_with_mqa.weight": "model-00136-of-000163.safetensors", + "model.layers.52.self_attn.kv_a_layernorm.weight": "model-00136-of-000163.safetensors", + "model.layers.52.self_attn.kv_b_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.self_attn.o_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.gate.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.gate.e_score_correction_bias": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.shared_experts.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.shared_experts.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.shared_experts.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.0.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.0.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.0.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.1.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.1.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.1.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.2.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.2.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.2.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.3.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.3.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.3.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.4.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.4.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.4.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.5.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.5.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.5.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.6.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.6.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.6.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.7.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.7.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.7.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.8.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.8.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.8.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.9.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.9.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.9.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.10.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.10.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.10.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.11.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.11.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.11.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.12.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.12.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.12.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.13.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.13.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.13.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.14.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.14.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.14.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.15.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.15.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.15.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.16.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.16.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.16.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.17.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.17.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.17.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.18.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.18.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.18.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.19.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.19.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.19.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.20.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.20.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.20.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.21.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.21.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.21.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.22.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.22.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.22.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.23.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.23.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.23.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.24.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.24.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.24.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.25.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.25.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.25.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.26.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.26.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.26.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.27.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.27.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.27.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.28.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.28.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.28.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.29.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.29.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.29.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.30.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.30.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.30.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.31.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.31.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.31.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.32.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.32.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.32.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.33.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.33.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.33.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.34.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.34.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.34.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.35.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.35.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.35.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.36.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.36.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.36.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.37.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.37.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.37.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.38.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.38.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.38.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.39.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.39.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.39.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.40.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.40.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.40.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.41.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.41.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.41.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.42.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.42.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.42.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.43.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.43.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.43.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.44.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.44.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.44.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.45.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.45.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.45.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.46.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.46.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.46.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.47.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.47.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.47.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.48.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.48.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.48.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.49.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.49.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.49.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.50.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.50.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.50.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.51.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.51.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.51.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.52.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.52.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.52.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.53.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.53.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.53.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.54.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.54.up_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.54.down_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.55.gate_proj.weight": "model-00136-of-000163.safetensors", + "model.layers.52.mlp.experts.55.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.55.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.56.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.56.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.56.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.57.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.57.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.57.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.58.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.58.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.58.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.59.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.59.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.59.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.60.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.60.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.60.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.61.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.61.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.61.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.62.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.62.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.62.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.63.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.63.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.63.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.64.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.64.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.64.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.65.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.65.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.65.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.66.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.66.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.66.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.67.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.67.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.67.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.68.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.68.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.68.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.69.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.69.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.69.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.70.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.70.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.70.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.71.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.71.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.71.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.72.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.72.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.72.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.73.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.73.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.73.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.74.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.74.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.74.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.75.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.75.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.75.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.76.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.76.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.76.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.77.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.77.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.77.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.78.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.78.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.78.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.79.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.79.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.79.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.80.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.80.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.80.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.81.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.81.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.81.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.82.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.82.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.82.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.83.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.83.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.83.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.84.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.84.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.84.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.85.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.85.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.85.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.86.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.86.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.86.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.87.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.87.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.87.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.88.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.88.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.88.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.89.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.89.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.89.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.90.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.90.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.90.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.91.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.91.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.91.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.92.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.92.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.92.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.93.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.93.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.93.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.94.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.94.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.94.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.95.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.95.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.95.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.96.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.96.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.96.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.97.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.97.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.97.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.98.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.98.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.98.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.99.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.99.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.99.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.100.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.100.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.100.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.101.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.101.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.101.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.102.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.102.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.102.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.103.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.103.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.103.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.104.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.104.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.104.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.105.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.105.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.105.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.106.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.106.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.106.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.107.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.107.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.107.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.108.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.108.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.108.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.109.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.109.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.109.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.110.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.110.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.110.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.111.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.111.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.111.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.112.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.112.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.112.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.113.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.113.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.113.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.114.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.114.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.114.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.115.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.115.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.115.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.116.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.116.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.116.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.117.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.117.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.117.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.118.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.118.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.118.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.119.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.119.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.119.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.120.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.120.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.120.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.121.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.121.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.121.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.122.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.122.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.122.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.123.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.123.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.123.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.124.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.124.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.124.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.125.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.125.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.125.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.126.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.126.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.126.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.127.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.127.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.127.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.128.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.128.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.128.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.129.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.129.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.129.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.130.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.130.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.130.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.131.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.131.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.131.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.132.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.132.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.132.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.133.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.133.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.133.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.134.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.134.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.134.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.135.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.135.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.135.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.136.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.136.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.136.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.137.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.137.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.137.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.138.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.138.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.138.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.139.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.139.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.139.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.140.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.140.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.140.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.141.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.141.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.141.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.142.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.142.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.142.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.143.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.143.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.143.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.144.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.144.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.144.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.145.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.145.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.145.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.146.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.146.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.146.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.147.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.147.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.147.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.148.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.148.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.148.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.149.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.149.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.149.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.150.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.150.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.150.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.151.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.151.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.151.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.152.gate_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.152.up_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.152.down_proj.weight": "model-00137-of-000163.safetensors", + "model.layers.52.mlp.experts.153.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.153.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.153.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.154.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.154.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.154.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.155.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.155.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.155.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.156.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.156.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.156.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.157.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.157.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.157.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.158.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.158.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.158.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.159.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.159.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.159.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.160.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.160.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.160.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.161.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.161.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.161.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.162.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.162.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.162.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.163.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.163.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.163.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.164.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.164.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.164.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.165.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.165.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.165.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.166.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.166.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.166.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.167.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.167.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.167.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.168.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.168.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.168.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.169.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.169.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.169.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.170.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.170.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.170.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.171.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.171.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.171.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.172.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.172.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.172.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.173.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.173.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.173.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.174.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.174.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.174.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.175.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.175.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.175.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.176.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.176.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.176.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.177.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.177.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.177.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.178.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.178.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.178.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.179.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.179.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.179.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.180.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.180.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.180.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.181.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.181.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.181.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.182.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.182.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.182.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.183.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.183.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.183.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.184.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.184.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.184.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.185.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.185.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.185.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.186.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.186.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.186.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.187.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.187.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.187.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.188.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.188.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.188.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.189.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.189.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.189.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.190.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.190.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.190.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.191.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.191.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.191.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.192.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.192.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.192.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.193.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.193.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.193.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.194.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.194.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.194.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.195.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.195.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.195.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.196.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.196.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.196.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.197.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.197.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.197.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.198.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.198.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.198.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.199.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.199.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.199.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.200.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.200.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.200.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.201.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.201.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.201.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.202.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.202.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.202.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.203.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.203.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.203.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.204.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.204.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.204.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.205.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.205.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.205.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.206.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.206.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.206.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.207.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.207.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.207.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.208.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.208.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.208.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.209.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.209.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.209.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.210.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.210.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.210.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.211.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.211.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.211.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.212.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.212.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.212.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.213.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.213.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.213.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.214.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.214.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.214.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.215.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.215.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.215.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.216.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.216.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.216.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.217.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.217.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.217.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.218.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.218.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.218.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.219.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.219.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.219.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.220.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.220.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.220.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.221.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.221.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.221.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.222.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.222.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.222.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.223.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.223.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.223.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.224.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.224.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.224.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.225.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.225.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.225.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.226.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.226.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.226.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.227.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.227.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.227.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.228.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.228.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.228.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.229.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.229.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.229.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.230.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.230.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.230.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.231.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.231.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.231.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.232.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.232.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.232.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.233.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.233.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.233.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.234.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.234.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.234.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.235.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.235.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.235.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.236.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.236.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.236.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.237.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.237.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.237.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.238.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.238.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.238.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.239.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.239.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.239.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.240.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.240.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.240.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.241.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.241.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.241.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.242.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.242.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.242.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.243.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.243.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.243.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.244.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.244.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.244.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.245.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.245.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.245.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.246.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.246.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.246.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.247.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.247.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.247.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.248.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.248.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.248.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.249.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.249.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.249.down_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.250.gate_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.250.up_proj.weight": "model-00138-of-000163.safetensors", + "model.layers.52.mlp.experts.250.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.52.mlp.experts.251.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.52.mlp.experts.251.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.52.mlp.experts.251.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.52.mlp.experts.252.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.52.mlp.experts.252.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.52.mlp.experts.252.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.52.mlp.experts.253.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.52.mlp.experts.253.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.52.mlp.experts.253.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.52.mlp.experts.254.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.52.mlp.experts.254.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.52.mlp.experts.254.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.52.mlp.experts.255.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.52.mlp.experts.255.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.52.mlp.experts.255.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.52.input_layernorm.weight": "model-00139-of-000163.safetensors", + "model.layers.52.post_attention_layernorm.weight": "model-00139-of-000163.safetensors", + "model.layers.53.self_attn.q_a_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.self_attn.q_a_layernorm.weight": "model-00139-of-000163.safetensors", + "model.layers.53.self_attn.q_b_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.self_attn.kv_a_proj_with_mqa.weight": "model-00139-of-000163.safetensors", + "model.layers.53.self_attn.kv_a_layernorm.weight": "model-00139-of-000163.safetensors", + "model.layers.53.self_attn.kv_b_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.self_attn.o_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.gate.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.gate.e_score_correction_bias": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.shared_experts.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.shared_experts.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.shared_experts.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.0.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.0.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.0.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.1.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.1.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.1.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.2.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.2.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.2.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.3.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.3.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.3.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.4.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.4.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.4.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.5.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.5.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.5.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.6.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.6.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.6.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.7.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.7.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.7.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.8.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.8.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.8.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.9.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.9.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.9.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.10.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.10.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.10.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.11.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.11.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.11.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.12.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.12.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.12.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.13.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.13.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.13.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.14.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.14.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.14.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.15.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.15.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.15.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.16.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.16.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.16.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.17.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.17.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.17.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.18.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.18.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.18.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.19.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.19.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.19.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.20.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.20.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.20.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.21.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.21.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.21.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.22.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.22.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.22.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.23.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.23.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.23.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.24.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.24.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.24.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.25.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.25.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.25.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.26.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.26.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.26.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.27.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.27.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.27.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.28.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.28.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.28.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.29.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.29.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.29.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.30.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.30.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.30.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.31.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.31.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.31.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.32.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.32.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.32.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.33.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.33.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.33.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.34.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.34.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.34.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.35.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.35.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.35.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.36.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.36.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.36.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.37.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.37.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.37.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.38.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.38.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.38.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.39.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.39.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.39.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.40.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.40.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.40.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.41.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.41.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.41.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.42.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.42.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.42.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.43.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.43.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.43.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.44.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.44.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.44.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.45.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.45.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.45.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.46.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.46.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.46.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.47.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.47.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.47.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.48.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.48.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.48.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.49.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.49.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.49.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.50.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.50.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.50.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.51.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.51.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.51.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.52.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.52.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.52.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.53.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.53.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.53.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.54.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.54.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.54.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.55.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.55.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.55.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.56.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.56.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.56.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.57.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.57.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.57.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.58.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.58.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.58.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.59.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.59.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.59.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.60.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.60.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.60.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.61.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.61.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.61.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.62.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.62.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.62.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.63.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.63.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.63.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.64.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.64.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.64.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.65.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.65.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.65.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.66.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.66.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.66.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.67.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.67.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.67.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.68.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.68.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.68.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.69.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.69.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.69.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.70.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.70.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.70.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.71.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.71.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.71.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.72.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.72.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.72.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.73.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.73.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.73.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.74.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.74.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.74.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.75.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.75.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.75.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.76.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.76.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.76.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.77.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.77.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.77.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.78.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.78.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.78.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.79.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.79.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.79.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.80.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.80.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.80.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.81.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.81.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.81.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.82.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.82.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.82.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.83.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.83.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.83.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.84.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.84.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.84.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.85.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.85.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.85.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.86.gate_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.86.up_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.86.down_proj.weight": "model-00139-of-000163.safetensors", + "model.layers.53.mlp.experts.87.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.87.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.87.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.88.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.88.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.88.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.89.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.89.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.89.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.90.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.90.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.90.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.91.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.91.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.91.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.92.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.92.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.92.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.93.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.93.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.93.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.94.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.94.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.94.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.95.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.95.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.95.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.96.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.96.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.96.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.97.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.97.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.97.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.98.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.98.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.98.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.99.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.99.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.99.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.100.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.100.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.100.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.101.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.101.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.101.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.102.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.102.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.102.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.103.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.103.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.103.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.104.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.104.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.104.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.105.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.105.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.105.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.106.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.106.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.106.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.107.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.107.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.107.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.108.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.108.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.108.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.109.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.109.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.109.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.110.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.110.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.110.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.111.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.111.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.111.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.112.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.112.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.112.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.113.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.113.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.113.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.114.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.114.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.114.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.115.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.115.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.115.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.116.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.116.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.116.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.117.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.117.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.117.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.118.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.118.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.118.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.119.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.119.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.119.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.120.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.120.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.120.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.121.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.121.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.121.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.122.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.122.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.122.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.123.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.123.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.123.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.124.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.124.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.124.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.125.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.125.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.125.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.126.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.126.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.126.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.127.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.127.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.127.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.128.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.128.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.128.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.129.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.129.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.129.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.130.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.130.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.130.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.131.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.131.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.131.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.132.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.132.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.132.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.133.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.133.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.133.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.134.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.134.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.134.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.135.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.135.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.135.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.136.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.136.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.136.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.137.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.137.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.137.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.138.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.138.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.138.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.139.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.139.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.139.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.140.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.140.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.140.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.141.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.141.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.141.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.142.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.142.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.142.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.143.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.143.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.143.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.144.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.144.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.144.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.145.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.145.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.145.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.146.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.146.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.146.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.147.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.147.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.147.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.148.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.148.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.148.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.149.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.149.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.149.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.150.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.150.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.150.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.151.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.151.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.151.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.152.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.152.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.152.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.153.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.153.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.153.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.154.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.154.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.154.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.155.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.155.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.155.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.156.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.156.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.156.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.157.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.157.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.157.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.158.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.158.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.158.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.159.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.159.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.159.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.160.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.160.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.160.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.161.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.161.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.161.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.162.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.162.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.162.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.163.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.163.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.163.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.164.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.164.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.164.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.165.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.165.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.165.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.166.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.166.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.166.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.167.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.167.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.167.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.168.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.168.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.168.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.169.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.169.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.169.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.170.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.170.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.170.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.171.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.171.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.171.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.172.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.172.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.172.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.173.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.173.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.173.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.174.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.174.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.174.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.175.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.175.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.175.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.176.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.176.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.176.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.177.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.177.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.177.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.178.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.178.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.178.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.179.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.179.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.179.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.180.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.180.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.180.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.181.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.181.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.181.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.182.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.182.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.182.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.183.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.183.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.183.down_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.184.gate_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.184.up_proj.weight": "model-00140-of-000163.safetensors", + "model.layers.53.mlp.experts.184.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.185.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.185.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.185.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.186.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.186.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.186.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.187.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.187.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.187.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.188.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.188.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.188.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.189.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.189.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.189.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.190.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.190.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.190.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.191.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.191.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.191.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.192.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.192.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.192.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.193.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.193.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.193.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.194.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.194.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.194.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.195.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.195.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.195.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.196.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.196.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.196.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.197.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.197.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.197.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.198.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.198.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.198.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.199.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.199.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.199.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.200.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.200.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.200.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.201.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.201.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.201.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.202.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.202.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.202.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.203.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.203.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.203.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.204.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.204.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.204.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.205.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.205.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.205.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.206.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.206.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.206.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.207.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.207.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.207.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.208.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.208.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.208.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.209.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.209.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.209.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.210.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.210.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.210.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.211.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.211.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.211.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.212.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.212.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.212.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.213.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.213.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.213.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.214.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.214.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.214.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.215.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.215.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.215.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.216.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.216.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.216.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.217.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.217.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.217.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.218.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.218.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.218.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.219.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.219.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.219.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.220.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.220.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.220.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.221.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.221.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.221.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.222.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.222.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.222.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.223.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.223.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.223.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.224.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.224.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.224.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.225.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.225.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.225.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.226.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.226.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.226.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.227.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.227.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.227.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.228.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.228.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.228.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.229.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.229.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.229.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.230.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.230.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.230.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.231.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.231.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.231.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.232.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.232.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.232.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.233.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.233.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.233.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.234.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.234.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.234.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.235.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.235.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.235.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.236.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.236.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.236.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.237.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.237.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.237.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.238.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.238.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.238.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.239.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.239.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.239.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.240.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.240.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.240.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.241.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.241.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.241.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.242.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.242.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.242.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.243.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.243.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.243.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.244.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.244.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.244.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.245.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.245.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.245.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.246.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.246.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.246.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.247.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.247.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.247.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.248.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.248.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.248.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.249.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.249.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.249.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.250.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.250.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.250.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.251.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.251.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.251.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.252.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.252.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.252.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.253.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.253.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.253.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.254.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.254.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.254.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.255.gate_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.255.up_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.mlp.experts.255.down_proj.weight": "model-00141-of-000163.safetensors", + "model.layers.53.input_layernorm.weight": "model-00141-of-000163.safetensors", + "model.layers.53.post_attention_layernorm.weight": "model-00141-of-000163.safetensors", + "model.layers.54.self_attn.q_a_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.self_attn.q_a_layernorm.weight": "model-00142-of-000163.safetensors", + "model.layers.54.self_attn.q_b_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.self_attn.kv_a_proj_with_mqa.weight": "model-00142-of-000163.safetensors", + "model.layers.54.self_attn.kv_a_layernorm.weight": "model-00142-of-000163.safetensors", + "model.layers.54.self_attn.kv_b_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.self_attn.o_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.gate.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.gate.e_score_correction_bias": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.shared_experts.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.shared_experts.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.shared_experts.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.0.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.0.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.0.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.1.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.1.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.1.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.2.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.2.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.2.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.3.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.3.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.3.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.4.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.4.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.4.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.5.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.5.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.5.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.6.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.6.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.6.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.7.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.7.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.7.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.8.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.8.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.8.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.9.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.9.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.9.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.10.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.10.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.10.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.11.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.11.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.11.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.12.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.12.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.12.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.13.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.13.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.13.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.14.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.14.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.14.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.15.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.15.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.15.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.16.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.16.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.16.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.17.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.17.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.17.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.18.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.18.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.18.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.19.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.19.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.19.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.20.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.20.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.20.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.21.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.21.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.21.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.22.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.22.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.22.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.23.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.23.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.23.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.24.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.24.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.24.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.25.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.25.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.25.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.26.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.26.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.26.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.27.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.27.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.27.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.28.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.28.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.28.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.29.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.29.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.29.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.30.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.30.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.30.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.31.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.31.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.31.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.32.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.32.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.32.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.33.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.33.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.33.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.34.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.34.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.34.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.35.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.35.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.35.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.36.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.36.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.36.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.37.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.37.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.37.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.38.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.38.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.38.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.39.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.39.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.39.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.40.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.40.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.40.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.41.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.41.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.41.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.42.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.42.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.42.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.43.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.43.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.43.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.44.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.44.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.44.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.45.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.45.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.45.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.46.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.46.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.46.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.47.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.47.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.47.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.48.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.48.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.48.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.49.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.49.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.49.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.50.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.50.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.50.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.51.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.51.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.51.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.52.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.52.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.52.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.53.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.53.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.53.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.54.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.54.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.54.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.55.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.55.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.55.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.56.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.56.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.56.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.57.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.57.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.57.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.58.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.58.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.58.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.59.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.59.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.59.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.60.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.60.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.60.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.61.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.61.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.61.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.62.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.62.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.62.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.63.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.63.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.63.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.64.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.64.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.64.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.65.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.65.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.65.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.66.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.66.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.66.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.67.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.67.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.67.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.68.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.68.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.68.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.69.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.69.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.69.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.70.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.70.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.70.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.71.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.71.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.71.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.72.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.72.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.72.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.73.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.73.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.73.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.74.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.74.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.74.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.75.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.75.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.75.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.76.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.76.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.76.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.77.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.77.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.77.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.78.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.78.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.78.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.79.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.79.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.79.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.80.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.80.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.80.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.81.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.81.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.81.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.82.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.82.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.82.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.83.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.83.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.83.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.84.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.84.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.84.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.85.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.85.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.85.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.86.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.86.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.86.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.87.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.87.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.87.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.88.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.88.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.88.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.89.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.89.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.89.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.90.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.90.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.90.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.91.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.91.up_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.91.down_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.92.gate_proj.weight": "model-00142-of-000163.safetensors", + "model.layers.54.mlp.experts.92.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.92.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.93.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.93.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.93.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.94.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.94.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.94.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.95.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.95.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.95.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.96.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.96.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.96.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.97.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.97.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.97.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.98.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.98.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.98.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.99.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.99.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.99.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.100.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.100.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.100.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.101.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.101.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.101.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.102.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.102.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.102.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.103.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.103.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.103.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.104.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.104.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.104.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.105.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.105.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.105.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.106.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.106.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.106.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.107.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.107.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.107.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.108.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.108.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.108.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.109.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.109.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.109.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.110.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.110.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.110.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.111.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.111.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.111.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.112.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.112.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.112.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.113.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.113.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.113.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.114.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.114.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.114.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.115.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.115.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.115.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.116.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.116.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.116.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.117.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.117.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.117.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.118.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.118.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.118.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.119.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.119.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.119.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.120.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.120.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.120.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.121.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.121.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.121.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.122.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.122.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.122.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.123.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.123.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.123.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.124.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.124.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.124.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.125.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.125.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.125.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.126.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.126.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.126.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.127.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.127.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.127.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.128.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.128.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.128.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.129.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.129.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.129.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.130.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.130.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.130.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.131.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.131.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.131.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.132.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.132.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.132.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.133.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.133.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.133.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.134.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.134.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.134.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.135.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.135.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.135.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.136.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.136.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.136.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.137.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.137.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.137.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.138.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.138.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.138.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.139.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.139.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.139.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.140.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.140.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.140.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.141.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.141.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.141.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.142.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.142.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.142.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.143.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.143.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.143.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.144.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.144.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.144.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.145.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.145.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.145.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.146.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.146.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.146.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.147.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.147.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.147.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.148.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.148.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.148.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.149.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.149.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.149.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.150.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.150.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.150.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.151.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.151.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.151.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.152.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.152.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.152.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.153.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.153.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.153.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.154.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.154.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.154.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.155.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.155.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.155.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.156.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.156.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.156.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.157.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.157.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.157.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.158.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.158.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.158.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.159.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.159.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.159.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.160.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.160.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.160.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.161.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.161.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.161.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.162.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.162.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.162.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.163.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.163.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.163.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.164.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.164.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.164.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.165.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.165.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.165.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.166.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.166.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.166.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.167.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.167.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.167.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.168.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.168.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.168.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.169.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.169.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.169.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.170.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.170.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.170.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.171.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.171.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.171.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.172.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.172.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.172.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.173.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.173.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.173.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.174.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.174.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.174.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.175.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.175.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.175.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.176.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.176.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.176.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.177.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.177.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.177.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.178.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.178.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.178.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.179.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.179.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.179.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.180.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.180.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.180.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.181.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.181.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.181.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.182.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.182.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.182.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.183.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.183.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.183.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.184.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.184.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.184.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.185.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.185.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.185.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.186.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.186.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.186.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.187.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.187.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.187.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.188.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.188.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.188.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.189.gate_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.189.up_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.189.down_proj.weight": "model-00143-of-000163.safetensors", + "model.layers.54.mlp.experts.190.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.190.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.190.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.191.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.191.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.191.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.192.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.192.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.192.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.193.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.193.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.193.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.194.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.194.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.194.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.195.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.195.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.195.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.196.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.196.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.196.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.197.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.197.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.197.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.198.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.198.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.198.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.199.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.199.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.199.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.200.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.200.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.200.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.201.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.201.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.201.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.202.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.202.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.202.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.203.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.203.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.203.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.204.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.204.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.204.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.205.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.205.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.205.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.206.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.206.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.206.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.207.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.207.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.207.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.208.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.208.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.208.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.209.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.209.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.209.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.210.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.210.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.210.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.211.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.211.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.211.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.212.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.212.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.212.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.213.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.213.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.213.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.214.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.214.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.214.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.215.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.215.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.215.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.216.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.216.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.216.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.217.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.217.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.217.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.218.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.218.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.218.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.219.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.219.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.219.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.220.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.220.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.220.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.221.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.221.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.221.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.222.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.222.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.222.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.223.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.223.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.223.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.224.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.224.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.224.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.225.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.225.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.225.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.226.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.226.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.226.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.227.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.227.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.227.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.228.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.228.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.228.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.229.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.229.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.229.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.230.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.230.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.230.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.231.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.231.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.231.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.232.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.232.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.232.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.233.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.233.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.233.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.234.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.234.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.234.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.235.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.235.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.235.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.236.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.236.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.236.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.237.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.237.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.237.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.238.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.238.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.238.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.239.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.239.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.239.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.240.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.240.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.240.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.241.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.241.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.241.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.242.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.242.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.242.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.243.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.243.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.243.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.244.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.244.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.244.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.245.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.245.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.245.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.246.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.246.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.246.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.247.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.247.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.247.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.248.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.248.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.248.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.249.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.249.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.249.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.250.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.250.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.250.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.251.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.251.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.251.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.252.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.252.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.252.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.253.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.253.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.253.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.254.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.254.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.254.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.255.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.255.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.mlp.experts.255.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.54.input_layernorm.weight": "model-00144-of-000163.safetensors", + "model.layers.54.post_attention_layernorm.weight": "model-00144-of-000163.safetensors", + "model.layers.55.self_attn.q_a_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.self_attn.q_a_layernorm.weight": "model-00144-of-000163.safetensors", + "model.layers.55.self_attn.q_b_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.self_attn.kv_a_proj_with_mqa.weight": "model-00144-of-000163.safetensors", + "model.layers.55.self_attn.kv_a_layernorm.weight": "model-00144-of-000163.safetensors", + "model.layers.55.self_attn.kv_b_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.self_attn.o_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.gate.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.gate.e_score_correction_bias": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.shared_experts.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.shared_experts.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.shared_experts.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.0.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.0.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.0.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.1.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.1.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.1.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.2.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.2.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.2.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.3.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.3.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.3.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.4.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.4.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.4.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.5.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.5.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.5.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.6.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.6.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.6.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.7.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.7.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.7.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.8.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.8.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.8.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.9.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.9.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.9.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.10.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.10.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.10.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.11.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.11.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.11.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.12.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.12.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.12.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.13.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.13.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.13.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.14.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.14.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.14.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.15.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.15.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.15.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.16.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.16.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.16.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.17.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.17.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.17.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.18.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.18.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.18.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.19.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.19.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.19.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.20.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.20.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.20.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.21.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.21.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.21.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.22.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.22.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.22.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.23.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.23.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.23.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.24.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.24.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.24.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.25.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.25.up_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.25.down_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.26.gate_proj.weight": "model-00144-of-000163.safetensors", + "model.layers.55.mlp.experts.26.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.26.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.27.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.27.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.27.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.28.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.28.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.28.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.29.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.29.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.29.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.30.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.30.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.30.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.31.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.31.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.31.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.32.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.32.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.32.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.33.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.33.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.33.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.34.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.34.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.34.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.35.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.35.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.35.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.36.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.36.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.36.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.37.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.37.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.37.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.38.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.38.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.38.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.39.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.39.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.39.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.40.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.40.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.40.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.41.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.41.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.41.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.42.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.42.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.42.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.43.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.43.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.43.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.44.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.44.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.44.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.45.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.45.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.45.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.46.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.46.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.46.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.47.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.47.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.47.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.48.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.48.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.48.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.49.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.49.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.49.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.50.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.50.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.50.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.51.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.51.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.51.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.52.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.52.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.52.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.53.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.53.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.53.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.54.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.54.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.54.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.55.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.55.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.55.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.56.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.56.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.56.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.57.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.57.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.57.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.58.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.58.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.58.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.59.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.59.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.59.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.60.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.60.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.60.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.61.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.61.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.61.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.62.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.62.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.62.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.63.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.63.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.63.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.64.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.64.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.64.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.65.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.65.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.65.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.66.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.66.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.66.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.67.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.67.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.67.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.68.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.68.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.68.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.69.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.69.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.69.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.70.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.70.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.70.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.71.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.71.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.71.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.72.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.72.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.72.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.73.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.73.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.73.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.74.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.74.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.74.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.75.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.75.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.75.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.76.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.76.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.76.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.77.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.77.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.77.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.78.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.78.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.78.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.79.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.79.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.79.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.80.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.80.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.80.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.81.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.81.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.81.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.82.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.82.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.82.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.83.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.83.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.83.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.84.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.84.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.84.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.85.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.85.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.85.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.86.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.86.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.86.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.87.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.87.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.87.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.88.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.88.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.88.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.89.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.89.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.89.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.90.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.90.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.90.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.91.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.91.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.91.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.92.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.92.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.92.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.93.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.93.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.93.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.94.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.94.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.94.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.95.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.95.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.95.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.96.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.96.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.96.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.97.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.97.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.97.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.98.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.98.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.98.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.99.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.99.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.99.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.100.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.100.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.100.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.101.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.101.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.101.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.102.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.102.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.102.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.103.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.103.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.103.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.104.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.104.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.104.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.105.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.105.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.105.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.106.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.106.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.106.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.107.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.107.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.107.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.108.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.108.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.108.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.109.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.109.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.109.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.110.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.110.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.110.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.111.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.111.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.111.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.112.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.112.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.112.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.113.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.113.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.113.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.114.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.114.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.114.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.115.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.115.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.115.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.116.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.116.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.116.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.117.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.117.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.117.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.118.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.118.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.118.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.119.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.119.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.119.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.120.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.120.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.120.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.121.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.121.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.121.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.122.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.122.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.122.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.123.gate_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.123.up_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.123.down_proj.weight": "model-00145-of-000163.safetensors", + "model.layers.55.mlp.experts.124.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.124.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.124.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.125.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.125.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.125.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.126.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.126.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.126.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.127.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.127.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.127.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.128.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.128.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.128.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.129.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.129.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.129.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.130.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.130.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.130.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.131.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.131.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.131.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.132.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.132.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.132.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.133.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.133.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.133.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.134.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.134.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.134.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.135.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.135.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.135.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.136.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.136.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.136.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.137.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.137.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.137.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.138.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.138.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.138.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.139.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.139.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.139.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.140.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.140.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.140.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.141.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.141.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.141.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.142.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.142.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.142.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.143.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.143.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.143.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.144.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.144.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.144.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.145.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.145.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.145.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.146.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.146.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.146.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.147.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.147.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.147.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.148.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.148.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.148.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.149.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.149.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.149.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.150.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.150.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.150.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.151.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.151.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.151.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.152.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.152.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.152.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.153.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.153.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.153.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.154.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.154.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.154.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.155.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.155.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.155.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.156.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.156.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.156.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.157.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.157.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.157.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.158.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.158.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.158.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.159.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.159.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.159.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.160.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.160.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.160.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.161.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.161.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.161.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.162.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.162.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.162.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.163.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.163.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.163.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.164.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.164.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.164.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.165.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.165.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.165.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.166.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.166.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.166.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.167.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.167.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.167.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.168.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.168.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.168.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.169.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.169.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.169.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.170.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.170.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.170.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.171.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.171.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.171.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.172.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.172.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.172.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.173.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.173.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.173.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.174.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.174.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.174.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.175.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.175.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.175.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.176.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.176.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.176.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.177.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.177.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.177.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.178.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.178.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.178.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.179.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.179.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.179.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.180.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.180.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.180.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.181.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.181.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.181.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.182.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.182.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.182.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.183.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.183.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.183.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.184.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.184.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.184.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.185.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.185.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.185.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.186.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.186.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.186.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.187.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.187.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.187.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.188.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.188.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.188.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.189.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.189.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.189.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.190.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.190.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.190.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.191.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.191.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.191.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.192.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.192.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.192.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.193.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.193.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.193.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.194.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.194.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.194.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.195.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.195.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.195.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.196.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.196.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.196.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.197.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.197.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.197.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.198.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.198.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.198.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.199.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.199.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.199.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.200.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.200.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.200.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.201.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.201.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.201.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.202.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.202.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.202.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.203.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.203.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.203.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.204.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.204.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.204.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.205.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.205.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.205.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.206.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.206.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.206.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.207.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.207.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.207.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.208.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.208.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.208.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.209.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.209.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.209.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.210.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.210.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.210.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.211.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.211.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.211.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.212.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.212.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.212.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.213.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.213.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.213.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.214.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.214.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.214.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.215.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.215.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.215.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.216.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.216.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.216.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.217.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.217.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.217.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.218.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.218.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.218.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.219.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.219.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.219.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.220.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.220.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.220.down_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.221.gate_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.221.up_proj.weight": "model-00146-of-000163.safetensors", + "model.layers.55.mlp.experts.221.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.222.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.222.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.222.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.223.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.223.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.223.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.224.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.224.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.224.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.225.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.225.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.225.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.226.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.226.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.226.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.227.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.227.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.227.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.228.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.228.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.228.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.229.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.229.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.229.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.230.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.230.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.230.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.231.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.231.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.231.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.232.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.232.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.232.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.233.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.233.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.233.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.234.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.234.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.234.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.235.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.235.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.235.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.236.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.236.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.236.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.237.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.237.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.237.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.238.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.238.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.238.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.239.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.239.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.239.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.240.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.240.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.240.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.241.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.241.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.241.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.242.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.242.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.242.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.243.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.243.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.243.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.244.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.244.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.244.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.245.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.245.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.245.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.246.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.246.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.246.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.247.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.247.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.247.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.248.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.248.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.248.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.249.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.249.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.249.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.250.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.250.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.250.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.251.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.251.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.251.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.252.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.252.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.252.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.253.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.253.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.253.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.254.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.254.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.254.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.255.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.255.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.mlp.experts.255.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.55.input_layernorm.weight": "model-00147-of-000163.safetensors", + "model.layers.55.post_attention_layernorm.weight": "model-00147-of-000163.safetensors", + "model.layers.56.self_attn.q_a_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.self_attn.q_a_layernorm.weight": "model-00147-of-000163.safetensors", + "model.layers.56.self_attn.q_b_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.self_attn.kv_a_proj_with_mqa.weight": "model-00147-of-000163.safetensors", + "model.layers.56.self_attn.kv_a_layernorm.weight": "model-00147-of-000163.safetensors", + "model.layers.56.self_attn.kv_b_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.self_attn.o_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.gate.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.gate.e_score_correction_bias": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.shared_experts.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.shared_experts.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.shared_experts.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.0.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.0.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.0.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.1.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.1.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.1.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.2.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.2.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.2.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.3.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.3.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.3.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.4.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.4.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.4.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.5.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.5.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.5.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.6.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.6.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.6.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.7.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.7.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.7.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.8.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.8.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.8.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.9.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.9.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.9.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.10.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.10.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.10.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.11.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.11.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.11.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.12.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.12.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.12.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.13.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.13.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.13.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.14.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.14.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.14.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.15.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.15.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.15.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.16.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.16.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.16.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.17.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.17.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.17.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.18.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.18.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.18.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.19.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.19.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.19.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.20.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.20.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.20.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.21.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.21.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.21.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.22.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.22.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.22.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.23.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.23.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.23.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.24.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.24.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.24.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.25.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.25.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.25.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.26.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.26.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.26.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.27.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.27.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.27.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.28.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.28.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.28.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.29.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.29.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.29.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.30.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.30.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.30.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.31.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.31.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.31.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.32.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.32.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.32.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.33.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.33.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.33.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.34.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.34.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.34.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.35.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.35.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.35.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.36.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.36.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.36.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.37.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.37.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.37.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.38.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.38.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.38.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.39.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.39.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.39.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.40.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.40.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.40.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.41.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.41.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.41.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.42.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.42.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.42.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.43.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.43.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.43.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.44.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.44.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.44.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.45.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.45.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.45.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.46.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.46.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.46.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.47.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.47.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.47.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.48.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.48.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.48.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.49.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.49.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.49.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.50.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.50.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.50.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.51.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.51.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.51.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.52.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.52.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.52.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.53.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.53.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.53.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.54.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.54.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.54.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.55.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.55.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.55.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.56.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.56.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.56.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.57.gate_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.57.up_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.57.down_proj.weight": "model-00147-of-000163.safetensors", + "model.layers.56.mlp.experts.58.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.58.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.58.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.59.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.59.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.59.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.60.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.60.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.60.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.61.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.61.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.61.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.62.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.62.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.62.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.63.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.63.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.63.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.64.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.64.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.64.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.65.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.65.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.65.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.66.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.66.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.66.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.67.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.67.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.67.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.68.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.68.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.68.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.69.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.69.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.69.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.70.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.70.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.70.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.71.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.71.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.71.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.72.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.72.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.72.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.73.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.73.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.73.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.74.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.74.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.74.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.75.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.75.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.75.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.76.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.76.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.76.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.77.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.77.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.77.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.78.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.78.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.78.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.79.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.79.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.79.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.80.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.80.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.80.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.81.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.81.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.81.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.82.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.82.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.82.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.83.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.83.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.83.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.84.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.84.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.84.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.85.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.85.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.85.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.86.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.86.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.86.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.87.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.87.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.87.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.88.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.88.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.88.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.89.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.89.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.89.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.90.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.90.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.90.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.91.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.91.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.91.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.92.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.92.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.92.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.93.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.93.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.93.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.94.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.94.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.94.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.95.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.95.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.95.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.96.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.96.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.96.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.97.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.97.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.97.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.98.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.98.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.98.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.99.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.99.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.99.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.100.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.100.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.100.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.101.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.101.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.101.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.102.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.102.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.102.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.103.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.103.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.103.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.104.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.104.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.104.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.105.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.105.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.105.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.106.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.106.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.106.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.107.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.107.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.107.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.108.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.108.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.108.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.109.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.109.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.109.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.110.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.110.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.110.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.111.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.111.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.111.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.112.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.112.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.112.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.113.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.113.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.113.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.114.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.114.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.114.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.115.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.115.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.115.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.116.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.116.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.116.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.117.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.117.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.117.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.118.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.118.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.118.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.119.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.119.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.119.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.120.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.120.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.120.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.121.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.121.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.121.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.122.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.122.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.122.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.123.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.123.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.123.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.124.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.124.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.124.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.125.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.125.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.125.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.126.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.126.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.126.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.127.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.127.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.127.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.128.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.128.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.128.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.129.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.129.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.129.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.130.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.130.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.130.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.131.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.131.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.131.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.132.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.132.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.132.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.133.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.133.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.133.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.134.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.134.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.134.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.135.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.135.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.135.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.136.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.136.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.136.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.137.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.137.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.137.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.138.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.138.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.138.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.139.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.139.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.139.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.140.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.140.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.140.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.141.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.141.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.141.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.142.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.142.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.142.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.143.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.143.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.143.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.144.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.144.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.144.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.145.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.145.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.145.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.146.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.146.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.146.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.147.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.147.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.147.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.148.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.148.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.148.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.149.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.149.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.149.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.150.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.150.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.150.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.151.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.151.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.151.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.152.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.152.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.152.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.153.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.153.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.153.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.154.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.154.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.154.down_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.155.gate_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.155.up_proj.weight": "model-00148-of-000163.safetensors", + "model.layers.56.mlp.experts.155.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.156.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.156.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.156.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.157.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.157.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.157.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.158.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.158.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.158.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.159.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.159.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.159.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.160.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.160.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.160.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.161.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.161.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.161.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.162.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.162.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.162.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.163.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.163.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.163.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.164.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.164.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.164.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.165.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.165.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.165.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.166.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.166.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.166.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.167.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.167.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.167.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.168.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.168.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.168.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.169.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.169.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.169.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.170.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.170.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.170.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.171.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.171.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.171.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.172.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.172.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.172.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.173.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.173.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.173.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.174.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.174.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.174.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.175.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.175.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.175.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.176.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.176.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.176.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.177.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.177.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.177.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.178.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.178.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.178.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.179.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.179.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.179.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.180.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.180.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.180.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.181.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.181.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.181.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.182.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.182.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.182.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.183.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.183.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.183.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.184.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.184.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.184.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.185.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.185.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.185.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.186.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.186.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.186.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.187.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.187.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.187.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.188.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.188.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.188.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.189.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.189.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.189.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.190.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.190.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.190.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.191.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.191.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.191.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.192.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.192.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.192.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.193.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.193.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.193.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.194.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.194.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.194.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.195.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.195.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.195.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.196.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.196.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.196.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.197.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.197.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.197.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.198.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.198.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.198.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.199.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.199.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.199.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.200.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.200.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.200.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.201.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.201.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.201.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.202.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.202.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.202.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.203.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.203.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.203.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.204.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.204.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.204.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.205.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.205.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.205.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.206.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.206.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.206.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.207.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.207.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.207.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.208.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.208.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.208.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.209.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.209.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.209.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.210.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.210.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.210.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.211.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.211.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.211.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.212.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.212.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.212.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.213.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.213.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.213.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.214.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.214.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.214.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.215.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.215.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.215.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.216.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.216.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.216.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.217.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.217.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.217.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.218.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.218.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.218.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.219.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.219.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.219.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.220.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.220.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.220.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.221.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.221.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.221.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.222.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.222.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.222.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.223.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.223.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.223.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.224.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.224.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.224.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.225.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.225.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.225.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.226.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.226.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.226.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.227.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.227.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.227.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.228.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.228.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.228.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.229.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.229.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.229.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.230.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.230.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.230.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.231.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.231.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.231.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.232.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.232.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.232.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.233.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.233.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.233.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.234.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.234.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.234.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.235.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.235.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.235.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.236.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.236.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.236.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.237.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.237.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.237.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.238.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.238.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.238.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.239.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.239.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.239.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.240.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.240.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.240.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.241.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.241.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.241.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.242.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.242.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.242.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.243.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.243.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.243.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.244.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.244.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.244.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.245.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.245.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.245.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.246.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.246.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.246.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.247.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.247.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.247.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.248.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.248.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.248.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.249.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.249.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.249.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.250.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.250.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.250.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.251.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.251.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.251.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.252.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.252.up_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.252.down_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.253.gate_proj.weight": "model-00149-of-000163.safetensors", + "model.layers.56.mlp.experts.253.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.56.mlp.experts.253.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.56.mlp.experts.254.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.56.mlp.experts.254.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.56.mlp.experts.254.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.56.mlp.experts.255.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.56.mlp.experts.255.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.56.mlp.experts.255.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.56.input_layernorm.weight": "model-00150-of-000163.safetensors", + "model.layers.56.post_attention_layernorm.weight": "model-00150-of-000163.safetensors", + "model.layers.57.self_attn.q_a_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.self_attn.q_a_layernorm.weight": "model-00150-of-000163.safetensors", + "model.layers.57.self_attn.q_b_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.self_attn.kv_a_proj_with_mqa.weight": "model-00150-of-000163.safetensors", + "model.layers.57.self_attn.kv_a_layernorm.weight": "model-00150-of-000163.safetensors", + "model.layers.57.self_attn.kv_b_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.self_attn.o_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.gate.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.gate.e_score_correction_bias": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.shared_experts.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.shared_experts.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.shared_experts.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.0.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.0.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.0.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.1.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.1.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.1.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.2.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.2.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.2.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.3.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.3.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.3.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.4.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.4.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.4.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.5.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.5.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.5.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.6.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.6.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.6.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.7.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.7.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.7.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.8.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.8.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.8.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.9.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.9.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.9.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.10.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.10.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.10.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.11.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.11.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.11.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.12.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.12.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.12.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.13.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.13.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.13.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.14.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.14.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.14.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.15.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.15.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.15.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.16.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.16.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.16.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.17.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.17.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.17.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.18.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.18.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.18.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.19.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.19.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.19.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.20.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.20.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.20.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.21.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.21.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.21.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.22.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.22.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.22.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.23.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.23.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.23.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.24.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.24.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.24.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.25.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.25.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.25.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.26.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.26.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.26.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.27.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.27.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.27.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.28.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.28.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.28.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.29.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.29.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.29.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.30.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.30.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.30.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.31.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.31.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.31.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.32.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.32.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.32.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.33.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.33.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.33.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.34.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.34.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.34.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.35.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.35.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.35.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.36.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.36.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.36.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.37.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.37.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.37.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.38.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.38.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.38.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.39.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.39.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.39.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.40.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.40.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.40.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.41.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.41.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.41.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.42.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.42.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.42.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.43.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.43.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.43.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.44.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.44.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.44.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.45.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.45.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.45.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.46.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.46.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.46.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.47.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.47.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.47.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.48.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.48.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.48.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.49.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.49.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.49.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.50.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.50.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.50.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.51.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.51.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.51.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.52.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.52.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.52.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.53.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.53.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.53.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.54.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.54.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.54.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.55.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.55.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.55.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.56.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.56.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.56.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.57.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.57.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.57.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.58.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.58.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.58.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.59.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.59.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.59.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.60.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.60.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.60.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.61.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.61.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.61.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.62.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.62.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.62.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.63.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.63.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.63.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.64.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.64.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.64.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.65.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.65.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.65.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.66.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.66.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.66.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.67.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.67.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.67.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.68.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.68.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.68.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.69.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.69.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.69.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.70.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.70.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.70.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.71.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.71.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.71.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.72.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.72.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.72.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.73.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.73.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.73.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.74.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.74.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.74.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.75.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.75.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.75.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.76.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.76.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.76.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.77.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.77.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.77.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.78.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.78.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.78.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.79.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.79.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.79.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.80.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.80.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.80.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.81.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.81.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.81.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.82.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.82.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.82.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.83.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.83.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.83.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.84.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.84.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.84.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.85.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.85.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.85.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.86.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.86.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.86.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.87.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.87.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.87.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.88.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.88.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.88.down_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.89.gate_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.89.up_proj.weight": "model-00150-of-000163.safetensors", + "model.layers.57.mlp.experts.89.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.90.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.90.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.90.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.91.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.91.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.91.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.92.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.92.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.92.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.93.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.93.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.93.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.94.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.94.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.94.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.95.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.95.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.95.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.96.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.96.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.96.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.97.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.97.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.97.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.98.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.98.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.98.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.99.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.99.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.99.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.100.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.100.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.100.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.101.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.101.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.101.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.102.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.102.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.102.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.103.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.103.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.103.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.104.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.104.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.104.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.105.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.105.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.105.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.106.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.106.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.106.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.107.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.107.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.107.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.108.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.108.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.108.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.109.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.109.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.109.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.110.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.110.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.110.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.111.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.111.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.111.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.112.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.112.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.112.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.113.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.113.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.113.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.114.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.114.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.114.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.115.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.115.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.115.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.116.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.116.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.116.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.117.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.117.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.117.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.118.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.118.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.118.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.119.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.119.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.119.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.120.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.120.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.120.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.121.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.121.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.121.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.122.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.122.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.122.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.123.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.123.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.123.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.124.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.124.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.124.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.125.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.125.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.125.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.126.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.126.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.126.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.127.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.127.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.127.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.128.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.128.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.128.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.129.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.129.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.129.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.130.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.130.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.130.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.131.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.131.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.131.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.132.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.132.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.132.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.133.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.133.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.133.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.134.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.134.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.134.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.135.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.135.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.135.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.136.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.136.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.136.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.137.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.137.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.137.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.138.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.138.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.138.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.139.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.139.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.139.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.140.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.140.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.140.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.141.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.141.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.141.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.142.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.142.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.142.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.143.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.143.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.143.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.144.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.144.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.144.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.145.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.145.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.145.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.146.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.146.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.146.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.147.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.147.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.147.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.148.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.148.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.148.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.149.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.149.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.149.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.150.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.150.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.150.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.151.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.151.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.151.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.152.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.152.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.152.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.153.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.153.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.153.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.154.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.154.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.154.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.155.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.155.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.155.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.156.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.156.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.156.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.157.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.157.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.157.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.158.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.158.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.158.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.159.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.159.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.159.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.160.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.160.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.160.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.161.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.161.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.161.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.162.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.162.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.162.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.163.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.163.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.163.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.164.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.164.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.164.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.165.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.165.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.165.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.166.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.166.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.166.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.167.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.167.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.167.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.168.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.168.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.168.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.169.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.169.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.169.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.170.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.170.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.170.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.171.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.171.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.171.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.172.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.172.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.172.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.173.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.173.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.173.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.174.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.174.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.174.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.175.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.175.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.175.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.176.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.176.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.176.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.177.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.177.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.177.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.178.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.178.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.178.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.179.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.179.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.179.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.180.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.180.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.180.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.181.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.181.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.181.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.182.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.182.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.182.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.183.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.183.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.183.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.184.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.184.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.184.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.185.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.185.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.185.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.186.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.186.up_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.186.down_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.187.gate_proj.weight": "model-00151-of-000163.safetensors", + "model.layers.57.mlp.experts.187.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.187.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.188.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.188.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.188.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.189.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.189.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.189.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.190.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.190.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.190.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.191.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.191.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.191.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.192.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.192.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.192.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.193.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.193.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.193.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.194.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.194.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.194.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.195.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.195.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.195.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.196.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.196.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.196.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.197.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.197.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.197.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.198.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.198.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.198.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.199.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.199.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.199.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.200.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.200.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.200.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.201.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.201.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.201.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.202.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.202.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.202.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.203.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.203.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.203.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.204.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.204.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.204.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.205.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.205.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.205.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.206.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.206.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.206.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.207.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.207.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.207.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.208.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.208.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.208.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.209.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.209.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.209.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.210.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.210.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.210.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.211.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.211.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.211.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.212.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.212.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.212.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.213.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.213.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.213.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.214.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.214.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.214.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.215.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.215.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.215.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.216.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.216.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.216.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.217.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.217.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.217.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.218.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.218.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.218.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.219.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.219.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.219.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.220.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.220.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.220.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.221.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.221.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.221.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.222.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.222.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.222.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.223.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.223.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.223.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.224.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.224.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.224.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.225.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.225.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.225.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.226.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.226.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.226.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.227.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.227.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.227.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.228.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.228.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.228.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.229.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.229.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.229.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.230.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.230.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.230.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.231.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.231.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.231.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.232.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.232.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.232.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.233.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.233.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.233.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.234.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.234.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.234.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.235.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.235.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.235.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.236.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.236.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.236.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.237.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.237.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.237.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.238.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.238.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.238.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.239.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.239.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.239.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.240.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.240.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.240.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.241.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.241.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.241.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.242.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.242.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.242.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.243.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.243.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.243.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.244.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.244.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.244.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.245.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.245.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.245.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.246.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.246.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.246.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.247.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.247.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.247.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.248.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.248.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.248.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.249.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.249.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.249.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.250.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.250.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.250.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.251.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.251.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.251.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.252.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.252.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.252.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.253.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.253.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.253.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.254.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.254.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.254.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.255.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.255.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.mlp.experts.255.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.57.input_layernorm.weight": "model-00152-of-000163.safetensors", + "model.layers.57.post_attention_layernorm.weight": "model-00152-of-000163.safetensors", + "model.layers.58.self_attn.q_a_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.self_attn.q_a_layernorm.weight": "model-00152-of-000163.safetensors", + "model.layers.58.self_attn.q_b_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.self_attn.kv_a_proj_with_mqa.weight": "model-00152-of-000163.safetensors", + "model.layers.58.self_attn.kv_a_layernorm.weight": "model-00152-of-000163.safetensors", + "model.layers.58.self_attn.kv_b_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.self_attn.o_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.gate.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.gate.e_score_correction_bias": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.shared_experts.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.shared_experts.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.shared_experts.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.0.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.0.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.0.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.1.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.1.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.1.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.2.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.2.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.2.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.3.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.3.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.3.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.4.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.4.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.4.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.5.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.5.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.5.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.6.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.6.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.6.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.7.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.7.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.7.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.8.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.8.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.8.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.9.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.9.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.9.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.10.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.10.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.10.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.11.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.11.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.11.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.12.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.12.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.12.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.13.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.13.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.13.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.14.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.14.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.14.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.15.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.15.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.15.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.16.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.16.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.16.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.17.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.17.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.17.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.18.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.18.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.18.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.19.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.19.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.19.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.20.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.20.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.20.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.21.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.21.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.21.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.22.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.22.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.22.down_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.23.gate_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.23.up_proj.weight": "model-00152-of-000163.safetensors", + "model.layers.58.mlp.experts.23.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.24.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.24.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.24.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.25.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.25.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.25.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.26.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.26.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.26.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.27.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.27.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.27.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.28.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.28.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.28.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.29.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.29.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.29.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.30.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.30.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.30.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.31.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.31.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.31.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.32.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.32.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.32.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.33.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.33.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.33.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.34.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.34.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.34.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.35.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.35.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.35.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.36.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.36.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.36.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.37.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.37.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.37.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.38.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.38.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.38.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.39.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.39.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.39.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.40.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.40.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.40.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.41.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.41.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.41.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.42.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.42.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.42.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.43.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.43.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.43.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.44.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.44.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.44.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.45.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.45.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.45.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.46.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.46.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.46.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.47.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.47.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.47.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.48.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.48.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.48.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.49.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.49.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.49.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.50.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.50.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.50.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.51.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.51.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.51.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.52.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.52.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.52.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.53.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.53.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.53.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.54.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.54.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.54.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.55.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.55.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.55.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.56.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.56.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.56.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.57.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.57.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.57.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.58.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.58.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.58.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.59.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.59.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.59.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.60.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.60.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.60.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.61.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.61.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.61.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.62.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.62.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.62.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.63.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.63.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.63.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.64.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.64.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.64.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.65.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.65.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.65.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.66.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.66.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.66.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.67.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.67.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.67.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.68.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.68.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.68.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.69.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.69.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.69.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.70.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.70.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.70.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.71.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.71.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.71.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.72.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.72.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.72.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.73.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.73.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.73.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.74.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.74.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.74.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.75.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.75.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.75.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.76.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.76.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.76.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.77.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.77.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.77.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.78.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.78.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.78.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.79.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.79.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.79.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.80.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.80.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.80.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.81.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.81.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.81.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.82.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.82.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.82.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.83.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.83.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.83.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.84.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.84.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.84.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.85.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.85.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.85.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.86.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.86.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.86.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.87.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.87.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.87.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.88.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.88.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.88.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.89.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.89.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.89.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.90.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.90.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.90.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.91.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.91.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.91.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.92.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.92.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.92.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.93.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.93.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.93.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.94.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.94.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.94.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.95.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.95.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.95.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.96.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.96.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.96.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.97.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.97.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.97.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.98.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.98.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.98.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.99.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.99.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.99.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.100.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.100.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.100.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.101.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.101.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.101.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.102.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.102.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.102.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.103.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.103.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.103.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.104.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.104.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.104.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.105.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.105.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.105.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.106.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.106.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.106.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.107.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.107.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.107.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.108.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.108.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.108.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.109.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.109.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.109.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.110.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.110.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.110.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.111.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.111.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.111.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.112.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.112.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.112.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.113.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.113.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.113.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.114.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.114.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.114.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.115.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.115.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.115.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.116.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.116.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.116.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.117.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.117.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.117.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.118.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.118.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.118.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.119.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.119.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.119.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.120.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.120.up_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.120.down_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.121.gate_proj.weight": "model-00153-of-000163.safetensors", + "model.layers.58.mlp.experts.121.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.121.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.122.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.122.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.122.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.123.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.123.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.123.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.124.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.124.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.124.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.125.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.125.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.125.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.126.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.126.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.126.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.127.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.127.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.127.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.128.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.128.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.128.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.129.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.129.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.129.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.130.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.130.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.130.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.131.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.131.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.131.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.132.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.132.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.132.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.133.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.133.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.133.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.134.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.134.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.134.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.135.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.135.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.135.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.136.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.136.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.136.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.137.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.137.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.137.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.138.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.138.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.138.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.139.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.139.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.139.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.140.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.140.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.140.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.141.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.141.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.141.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.142.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.142.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.142.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.143.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.143.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.143.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.144.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.144.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.144.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.145.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.145.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.145.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.146.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.146.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.146.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.147.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.147.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.147.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.148.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.148.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.148.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.149.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.149.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.149.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.150.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.150.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.150.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.151.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.151.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.151.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.152.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.152.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.152.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.153.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.153.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.153.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.154.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.154.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.154.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.155.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.155.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.155.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.156.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.156.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.156.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.157.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.157.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.157.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.158.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.158.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.158.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.159.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.159.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.159.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.160.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.160.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.160.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.161.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.161.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.161.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.162.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.162.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.162.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.163.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.163.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.163.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.164.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.164.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.164.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.165.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.165.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.165.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.166.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.166.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.166.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.167.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.167.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.167.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.168.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.168.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.168.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.169.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.169.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.169.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.170.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.170.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.170.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.171.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.171.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.171.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.172.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.172.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.172.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.173.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.173.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.173.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.174.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.174.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.174.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.175.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.175.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.175.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.176.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.176.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.176.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.177.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.177.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.177.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.178.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.178.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.178.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.179.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.179.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.179.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.180.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.180.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.180.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.181.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.181.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.181.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.182.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.182.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.182.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.183.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.183.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.183.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.184.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.184.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.184.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.185.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.185.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.185.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.186.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.186.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.186.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.187.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.187.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.187.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.188.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.188.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.188.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.189.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.189.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.189.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.190.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.190.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.190.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.191.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.191.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.191.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.192.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.192.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.192.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.193.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.193.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.193.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.194.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.194.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.194.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.195.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.195.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.195.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.196.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.196.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.196.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.197.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.197.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.197.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.198.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.198.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.198.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.199.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.199.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.199.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.200.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.200.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.200.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.201.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.201.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.201.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.202.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.202.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.202.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.203.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.203.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.203.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.204.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.204.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.204.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.205.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.205.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.205.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.206.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.206.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.206.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.207.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.207.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.207.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.208.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.208.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.208.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.209.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.209.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.209.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.210.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.210.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.210.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.211.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.211.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.211.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.212.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.212.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.212.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.213.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.213.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.213.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.214.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.214.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.214.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.215.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.215.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.215.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.216.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.216.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.216.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.217.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.217.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.217.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.218.gate_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.218.up_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.218.down_proj.weight": "model-00154-of-000163.safetensors", + "model.layers.58.mlp.experts.219.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.219.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.219.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.220.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.220.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.220.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.221.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.221.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.221.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.222.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.222.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.222.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.223.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.223.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.223.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.224.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.224.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.224.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.225.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.225.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.225.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.226.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.226.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.226.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.227.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.227.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.227.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.228.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.228.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.228.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.229.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.229.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.229.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.230.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.230.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.230.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.231.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.231.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.231.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.232.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.232.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.232.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.233.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.233.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.233.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.234.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.234.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.234.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.235.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.235.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.235.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.236.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.236.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.236.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.237.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.237.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.237.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.238.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.238.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.238.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.239.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.239.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.239.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.240.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.240.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.240.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.241.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.241.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.241.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.242.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.242.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.242.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.243.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.243.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.243.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.244.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.244.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.244.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.245.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.245.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.245.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.246.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.246.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.246.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.247.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.247.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.247.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.248.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.248.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.248.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.249.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.249.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.249.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.250.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.250.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.250.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.251.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.251.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.251.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.252.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.252.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.252.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.253.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.253.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.253.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.254.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.254.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.254.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.255.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.255.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.mlp.experts.255.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.58.input_layernorm.weight": "model-00155-of-000163.safetensors", + "model.layers.58.post_attention_layernorm.weight": "model-00155-of-000163.safetensors", + "model.layers.59.self_attn.q_a_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.self_attn.q_a_layernorm.weight": "model-00155-of-000163.safetensors", + "model.layers.59.self_attn.q_b_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.self_attn.kv_a_proj_with_mqa.weight": "model-00155-of-000163.safetensors", + "model.layers.59.self_attn.kv_a_layernorm.weight": "model-00155-of-000163.safetensors", + "model.layers.59.self_attn.kv_b_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.self_attn.o_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.gate.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.gate.e_score_correction_bias": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.shared_experts.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.shared_experts.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.shared_experts.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.0.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.0.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.0.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.1.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.1.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.1.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.2.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.2.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.2.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.3.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.3.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.3.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.4.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.4.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.4.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.5.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.5.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.5.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.6.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.6.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.6.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.7.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.7.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.7.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.8.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.8.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.8.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.9.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.9.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.9.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.10.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.10.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.10.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.11.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.11.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.11.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.12.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.12.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.12.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.13.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.13.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.13.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.14.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.14.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.14.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.15.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.15.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.15.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.16.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.16.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.16.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.17.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.17.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.17.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.18.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.18.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.18.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.19.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.19.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.19.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.20.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.20.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.20.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.21.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.21.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.21.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.22.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.22.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.22.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.23.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.23.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.23.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.24.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.24.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.24.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.25.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.25.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.25.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.26.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.26.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.26.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.27.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.27.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.27.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.28.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.28.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.28.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.29.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.29.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.29.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.30.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.30.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.30.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.31.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.31.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.31.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.32.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.32.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.32.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.33.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.33.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.33.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.34.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.34.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.34.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.35.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.35.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.35.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.36.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.36.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.36.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.37.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.37.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.37.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.38.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.38.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.38.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.39.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.39.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.39.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.40.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.40.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.40.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.41.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.41.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.41.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.42.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.42.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.42.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.43.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.43.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.43.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.44.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.44.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.44.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.45.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.45.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.45.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.46.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.46.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.46.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.47.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.47.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.47.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.48.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.48.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.48.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.49.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.49.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.49.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.50.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.50.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.50.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.51.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.51.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.51.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.52.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.52.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.52.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.53.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.53.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.53.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.54.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.54.up_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.54.down_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.55.gate_proj.weight": "model-00155-of-000163.safetensors", + "model.layers.59.mlp.experts.55.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.55.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.56.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.56.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.56.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.57.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.57.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.57.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.58.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.58.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.58.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.59.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.59.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.59.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.60.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.60.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.60.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.61.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.61.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.61.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.62.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.62.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.62.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.63.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.63.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.63.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.64.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.64.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.64.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.65.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.65.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.65.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.66.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.66.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.66.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.67.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.67.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.67.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.68.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.68.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.68.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.69.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.69.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.69.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.70.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.70.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.70.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.71.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.71.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.71.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.72.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.72.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.72.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.73.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.73.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.73.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.74.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.74.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.74.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.75.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.75.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.75.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.76.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.76.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.76.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.77.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.77.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.77.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.78.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.78.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.78.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.79.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.79.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.79.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.80.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.80.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.80.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.81.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.81.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.81.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.82.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.82.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.82.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.83.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.83.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.83.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.84.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.84.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.84.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.85.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.85.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.85.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.86.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.86.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.86.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.87.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.87.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.87.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.88.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.88.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.88.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.89.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.89.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.89.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.90.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.90.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.90.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.91.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.91.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.91.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.92.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.92.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.92.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.93.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.93.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.93.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.94.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.94.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.94.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.95.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.95.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.95.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.96.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.96.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.96.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.97.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.97.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.97.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.98.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.98.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.98.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.99.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.99.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.99.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.100.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.100.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.100.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.101.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.101.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.101.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.102.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.102.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.102.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.103.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.103.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.103.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.104.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.104.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.104.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.105.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.105.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.105.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.106.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.106.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.106.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.107.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.107.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.107.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.108.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.108.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.108.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.109.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.109.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.109.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.110.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.110.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.110.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.111.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.111.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.111.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.112.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.112.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.112.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.113.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.113.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.113.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.114.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.114.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.114.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.115.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.115.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.115.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.116.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.116.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.116.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.117.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.117.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.117.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.118.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.118.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.118.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.119.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.119.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.119.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.120.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.120.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.120.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.121.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.121.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.121.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.122.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.122.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.122.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.123.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.123.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.123.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.124.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.124.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.124.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.125.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.125.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.125.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.126.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.126.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.126.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.127.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.127.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.127.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.128.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.128.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.128.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.129.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.129.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.129.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.130.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.130.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.130.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.131.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.131.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.131.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.132.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.132.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.132.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.133.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.133.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.133.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.134.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.134.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.134.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.135.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.135.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.135.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.136.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.136.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.136.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.137.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.137.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.137.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.138.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.138.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.138.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.139.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.139.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.139.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.140.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.140.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.140.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.141.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.141.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.141.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.142.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.142.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.142.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.143.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.143.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.143.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.144.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.144.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.144.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.145.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.145.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.145.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.146.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.146.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.146.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.147.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.147.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.147.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.148.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.148.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.148.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.149.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.149.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.149.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.150.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.150.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.150.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.151.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.151.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.151.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.152.gate_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.152.up_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.152.down_proj.weight": "model-00156-of-000163.safetensors", + "model.layers.59.mlp.experts.153.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.153.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.153.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.154.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.154.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.154.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.155.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.155.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.155.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.156.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.156.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.156.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.157.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.157.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.157.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.158.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.158.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.158.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.159.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.159.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.159.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.160.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.160.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.160.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.161.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.161.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.161.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.162.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.162.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.162.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.163.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.163.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.163.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.164.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.164.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.164.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.165.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.165.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.165.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.166.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.166.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.166.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.167.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.167.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.167.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.168.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.168.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.168.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.169.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.169.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.169.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.170.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.170.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.170.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.171.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.171.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.171.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.172.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.172.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.172.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.173.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.173.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.173.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.174.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.174.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.174.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.175.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.175.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.175.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.176.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.176.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.176.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.177.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.177.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.177.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.178.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.178.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.178.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.179.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.179.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.179.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.180.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.180.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.180.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.181.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.181.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.181.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.182.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.182.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.182.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.183.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.183.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.183.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.184.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.184.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.184.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.185.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.185.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.185.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.186.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.186.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.186.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.187.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.187.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.187.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.188.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.188.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.188.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.189.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.189.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.189.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.190.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.190.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.190.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.191.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.191.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.191.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.192.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.192.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.192.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.193.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.193.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.193.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.194.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.194.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.194.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.195.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.195.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.195.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.196.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.196.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.196.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.197.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.197.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.197.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.198.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.198.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.198.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.199.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.199.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.199.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.200.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.200.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.200.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.201.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.201.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.201.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.202.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.202.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.202.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.203.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.203.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.203.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.204.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.204.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.204.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.205.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.205.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.205.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.206.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.206.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.206.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.207.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.207.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.207.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.208.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.208.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.208.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.209.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.209.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.209.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.210.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.210.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.210.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.211.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.211.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.211.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.212.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.212.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.212.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.213.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.213.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.213.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.214.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.214.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.214.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.215.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.215.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.215.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.216.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.216.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.216.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.217.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.217.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.217.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.218.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.218.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.218.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.219.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.219.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.219.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.220.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.220.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.220.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.221.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.221.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.221.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.222.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.222.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.222.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.223.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.223.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.223.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.224.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.224.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.224.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.225.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.225.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.225.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.226.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.226.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.226.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.227.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.227.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.227.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.228.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.228.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.228.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.229.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.229.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.229.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.230.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.230.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.230.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.231.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.231.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.231.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.232.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.232.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.232.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.233.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.233.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.233.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.234.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.234.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.234.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.235.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.235.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.235.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.236.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.236.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.236.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.237.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.237.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.237.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.238.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.238.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.238.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.239.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.239.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.239.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.240.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.240.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.240.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.241.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.241.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.241.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.242.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.242.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.242.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.243.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.243.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.243.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.244.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.244.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.244.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.245.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.245.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.245.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.246.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.246.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.246.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.247.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.247.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.247.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.248.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.248.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.248.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.249.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.249.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.249.down_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.250.gate_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.250.up_proj.weight": "model-00157-of-000163.safetensors", + "model.layers.59.mlp.experts.250.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.59.mlp.experts.251.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.59.mlp.experts.251.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.59.mlp.experts.251.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.59.mlp.experts.252.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.59.mlp.experts.252.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.59.mlp.experts.252.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.59.mlp.experts.253.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.59.mlp.experts.253.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.59.mlp.experts.253.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.59.mlp.experts.254.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.59.mlp.experts.254.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.59.mlp.experts.254.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.59.mlp.experts.255.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.59.mlp.experts.255.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.59.mlp.experts.255.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.59.input_layernorm.weight": "model-00158-of-000163.safetensors", + "model.layers.59.post_attention_layernorm.weight": "model-00158-of-000163.safetensors", + "model.layers.60.self_attn.q_a_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.self_attn.q_a_layernorm.weight": "model-00158-of-000163.safetensors", + "model.layers.60.self_attn.q_b_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.self_attn.kv_a_proj_with_mqa.weight": "model-00158-of-000163.safetensors", + "model.layers.60.self_attn.kv_a_layernorm.weight": "model-00158-of-000163.safetensors", + "model.layers.60.self_attn.kv_b_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.self_attn.o_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.gate.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.gate.e_score_correction_bias": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.shared_experts.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.shared_experts.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.shared_experts.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.0.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.0.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.0.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.1.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.1.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.1.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.2.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.2.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.2.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.3.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.3.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.3.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.4.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.4.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.4.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.5.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.5.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.5.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.6.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.6.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.6.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.7.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.7.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.7.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.8.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.8.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.8.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.9.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.9.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.9.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.10.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.10.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.10.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.11.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.11.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.11.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.12.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.12.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.12.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.13.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.13.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.13.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.14.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.14.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.14.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.15.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.15.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.15.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.16.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.16.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.16.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.17.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.17.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.17.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.18.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.18.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.18.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.19.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.19.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.19.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.20.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.20.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.20.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.21.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.21.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.21.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.22.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.22.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.22.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.23.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.23.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.23.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.24.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.24.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.24.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.25.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.25.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.25.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.26.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.26.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.26.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.27.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.27.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.27.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.28.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.28.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.28.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.29.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.29.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.29.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.30.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.30.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.30.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.31.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.31.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.31.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.32.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.32.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.32.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.33.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.33.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.33.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.34.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.34.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.34.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.35.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.35.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.35.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.36.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.36.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.36.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.37.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.37.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.37.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.38.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.38.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.38.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.39.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.39.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.39.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.40.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.40.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.40.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.41.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.41.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.41.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.42.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.42.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.42.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.43.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.43.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.43.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.44.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.44.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.44.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.45.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.45.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.45.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.46.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.46.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.46.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.47.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.47.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.47.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.48.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.48.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.48.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.49.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.49.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.49.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.50.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.50.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.50.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.51.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.51.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.51.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.52.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.52.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.52.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.53.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.53.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.53.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.54.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.54.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.54.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.55.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.55.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.55.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.56.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.56.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.56.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.57.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.57.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.57.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.58.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.58.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.58.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.59.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.59.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.59.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.60.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.60.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.60.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.61.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.61.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.61.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.62.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.62.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.62.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.63.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.63.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.63.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.64.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.64.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.64.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.65.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.65.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.65.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.66.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.66.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.66.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.67.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.67.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.67.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.68.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.68.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.68.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.69.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.69.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.69.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.70.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.70.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.70.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.71.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.71.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.71.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.72.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.72.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.72.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.73.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.73.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.73.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.74.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.74.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.74.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.75.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.75.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.75.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.76.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.76.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.76.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.77.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.77.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.77.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.78.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.78.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.78.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.79.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.79.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.79.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.80.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.80.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.80.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.81.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.81.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.81.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.82.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.82.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.82.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.83.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.83.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.83.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.84.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.84.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.84.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.85.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.85.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.85.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.86.gate_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.86.up_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.86.down_proj.weight": "model-00158-of-000163.safetensors", + "model.layers.60.mlp.experts.87.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.87.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.87.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.88.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.88.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.88.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.89.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.89.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.89.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.90.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.90.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.90.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.91.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.91.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.91.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.92.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.92.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.92.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.93.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.93.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.93.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.94.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.94.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.94.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.95.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.95.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.95.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.96.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.96.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.96.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.97.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.97.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.97.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.98.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.98.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.98.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.99.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.99.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.99.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.100.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.100.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.100.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.101.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.101.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.101.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.102.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.102.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.102.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.103.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.103.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.103.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.104.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.104.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.104.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.105.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.105.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.105.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.106.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.106.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.106.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.107.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.107.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.107.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.108.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.108.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.108.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.109.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.109.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.109.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.110.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.110.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.110.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.111.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.111.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.111.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.112.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.112.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.112.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.113.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.113.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.113.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.114.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.114.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.114.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.115.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.115.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.115.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.116.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.116.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.116.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.117.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.117.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.117.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.118.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.118.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.118.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.119.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.119.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.119.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.120.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.120.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.120.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.121.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.121.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.121.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.122.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.122.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.122.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.123.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.123.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.123.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.124.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.124.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.124.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.125.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.125.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.125.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.126.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.126.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.126.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.127.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.127.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.127.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.128.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.128.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.128.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.129.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.129.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.129.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.130.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.130.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.130.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.131.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.131.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.131.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.132.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.132.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.132.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.133.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.133.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.133.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.134.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.134.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.134.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.135.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.135.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.135.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.136.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.136.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.136.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.137.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.137.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.137.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.138.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.138.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.138.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.139.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.139.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.139.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.140.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.140.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.140.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.141.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.141.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.141.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.142.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.142.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.142.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.143.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.143.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.143.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.144.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.144.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.144.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.145.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.145.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.145.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.146.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.146.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.146.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.147.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.147.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.147.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.148.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.148.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.148.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.149.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.149.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.149.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.150.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.150.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.150.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.151.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.151.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.151.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.152.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.152.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.152.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.153.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.153.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.153.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.154.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.154.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.154.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.155.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.155.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.155.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.156.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.156.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.156.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.157.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.157.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.157.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.158.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.158.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.158.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.159.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.159.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.159.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.160.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.160.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.160.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.161.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.161.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.161.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.162.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.162.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.162.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.163.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.163.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.163.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.164.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.164.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.164.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.165.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.165.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.165.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.166.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.166.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.166.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.167.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.167.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.167.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.168.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.168.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.168.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.169.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.169.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.169.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.170.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.170.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.170.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.171.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.171.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.171.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.172.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.172.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.172.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.173.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.173.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.173.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.174.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.174.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.174.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.175.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.175.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.175.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.176.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.176.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.176.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.177.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.177.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.177.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.178.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.178.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.178.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.179.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.179.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.179.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.180.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.180.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.180.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.181.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.181.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.181.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.182.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.182.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.182.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.183.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.183.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.183.down_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.184.gate_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.184.up_proj.weight": "model-00159-of-000163.safetensors", + "model.layers.60.mlp.experts.184.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.185.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.185.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.185.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.186.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.186.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.186.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.187.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.187.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.187.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.188.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.188.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.188.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.189.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.189.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.189.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.190.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.190.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.190.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.191.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.191.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.191.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.192.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.192.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.192.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.193.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.193.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.193.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.194.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.194.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.194.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.195.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.195.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.195.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.196.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.196.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.196.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.197.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.197.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.197.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.198.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.198.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.198.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.199.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.199.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.199.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.200.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.200.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.200.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.201.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.201.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.201.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.202.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.202.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.202.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.203.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.203.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.203.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.204.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.204.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.204.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.205.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.205.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.205.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.206.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.206.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.206.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.207.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.207.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.207.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.208.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.208.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.208.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.209.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.209.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.209.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.210.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.210.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.210.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.211.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.211.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.211.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.212.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.212.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.212.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.213.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.213.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.213.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.214.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.214.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.214.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.215.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.215.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.215.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.216.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.216.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.216.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.217.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.217.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.217.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.218.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.218.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.218.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.219.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.219.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.219.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.220.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.220.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.220.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.221.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.221.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.221.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.222.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.222.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.222.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.223.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.223.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.223.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.224.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.224.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.224.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.225.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.225.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.225.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.226.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.226.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.226.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.227.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.227.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.227.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.228.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.228.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.228.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.229.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.229.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.229.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.230.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.230.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.230.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.231.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.231.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.231.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.232.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.232.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.232.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.233.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.233.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.233.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.234.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.234.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.234.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.235.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.235.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.235.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.236.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.236.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.236.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.237.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.237.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.237.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.238.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.238.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.238.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.239.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.239.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.239.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.240.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.240.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.240.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.241.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.241.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.241.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.242.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.242.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.242.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.243.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.243.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.243.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.244.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.244.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.244.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.245.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.245.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.245.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.246.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.246.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.246.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.247.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.247.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.247.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.248.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.248.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.248.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.249.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.249.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.249.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.250.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.250.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.250.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.251.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.251.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.251.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.252.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.252.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.252.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.253.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.253.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.253.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.254.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.254.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.254.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.255.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.255.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.mlp.experts.255.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.60.input_layernorm.weight": "model-00160-of-000163.safetensors", + "model.layers.60.post_attention_layernorm.weight": "model-00160-of-000163.safetensors", + "model.norm.weight": "model-00160-of-000163.safetensors", + "lm_head.weight": "model-00160-of-000163.safetensors", + "model.layers.61.self_attn.q_a_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.61.self_attn.q_a_layernorm.weight": "model-00160-of-000163.safetensors", + "model.layers.61.self_attn.q_b_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.61.self_attn.kv_a_proj_with_mqa.weight": "model-00160-of-000163.safetensors", + "model.layers.61.self_attn.kv_a_layernorm.weight": "model-00160-of-000163.safetensors", + "model.layers.61.self_attn.kv_b_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.61.self_attn.o_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.61.mlp.gate.weight": "model-00160-of-000163.safetensors", + "model.layers.61.mlp.gate.e_score_correction_bias": "model-00160-of-000163.safetensors", + "model.layers.61.mlp.shared_experts.gate_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.61.mlp.shared_experts.up_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.61.mlp.shared_experts.down_proj.weight": "model-00160-of-000163.safetensors", + "model.layers.61.mlp.experts.0.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.0.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.0.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.1.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.1.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.1.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.2.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.2.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.2.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.3.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.3.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.3.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.4.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.4.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.4.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.5.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.5.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.5.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.6.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.6.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.6.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.7.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.7.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.7.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.8.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.8.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.8.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.9.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.9.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.9.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.10.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.10.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.10.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.11.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.11.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.11.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.12.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.12.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.12.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.13.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.13.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.13.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.14.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.14.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.14.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.15.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.15.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.15.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.16.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.16.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.16.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.17.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.17.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.17.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.18.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.18.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.18.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.19.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.19.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.19.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.20.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.20.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.20.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.21.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.21.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.21.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.22.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.22.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.22.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.23.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.23.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.23.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.24.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.24.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.24.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.25.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.25.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.25.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.26.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.26.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.26.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.27.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.27.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.27.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.28.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.28.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.28.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.29.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.29.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.29.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.30.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.30.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.30.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.31.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.31.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.31.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.32.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.32.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.32.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.33.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.33.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.33.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.34.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.34.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.34.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.35.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.35.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.35.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.36.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.36.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.36.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.37.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.37.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.37.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.38.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.38.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.38.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.39.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.39.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.39.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.40.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.40.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.40.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.41.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.41.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.41.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.42.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.42.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.42.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.43.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.43.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.43.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.44.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.44.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.44.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.45.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.45.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.45.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.46.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.46.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.46.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.47.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.47.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.47.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.48.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.48.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.48.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.49.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.49.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.49.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.50.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.50.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.50.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.51.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.51.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.51.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.52.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.52.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.52.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.53.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.53.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.53.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.54.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.54.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.54.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.55.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.55.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.55.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.56.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.56.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.56.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.57.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.57.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.57.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.58.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.58.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.58.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.59.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.59.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.59.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.60.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.60.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.60.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.61.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.61.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.61.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.62.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.62.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.62.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.63.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.63.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.63.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.64.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.64.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.64.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.65.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.65.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.65.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.66.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.66.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.66.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.67.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.67.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.67.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.68.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.68.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.68.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.69.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.69.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.69.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.70.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.70.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.70.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.71.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.71.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.71.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.72.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.72.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.72.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.73.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.73.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.73.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.74.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.74.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.74.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.75.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.75.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.75.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.76.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.76.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.76.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.77.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.77.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.77.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.78.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.78.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.78.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.79.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.79.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.79.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.80.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.80.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.80.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.81.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.81.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.81.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.82.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.82.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.82.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.83.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.83.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.83.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.84.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.84.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.84.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.85.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.85.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.85.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.86.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.86.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.86.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.87.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.87.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.87.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.88.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.88.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.88.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.89.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.89.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.89.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.90.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.90.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.90.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.91.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.91.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.91.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.92.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.92.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.92.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.93.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.93.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.93.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.94.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.94.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.94.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.95.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.95.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.95.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.96.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.96.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.96.down_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.97.gate_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.97.up_proj.weight": "model-00161-of-000163.safetensors", + "model.layers.61.mlp.experts.97.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.98.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.98.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.98.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.99.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.99.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.99.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.100.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.100.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.100.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.101.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.101.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.101.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.102.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.102.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.102.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.103.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.103.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.103.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.104.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.104.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.104.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.105.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.105.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.105.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.106.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.106.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.106.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.107.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.107.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.107.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.108.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.108.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.108.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.109.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.109.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.109.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.110.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.110.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.110.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.111.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.111.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.111.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.112.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.112.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.112.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.113.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.113.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.113.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.114.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.114.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.114.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.115.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.115.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.115.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.116.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.116.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.116.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.117.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.117.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.117.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.118.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.118.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.118.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.119.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.119.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.119.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.120.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.120.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.120.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.121.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.121.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.121.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.122.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.122.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.122.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.123.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.123.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.123.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.124.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.124.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.124.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.125.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.125.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.125.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.126.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.126.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.126.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.127.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.127.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.127.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.128.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.128.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.128.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.129.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.129.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.129.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.130.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.130.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.130.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.131.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.131.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.131.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.132.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.132.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.132.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.133.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.133.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.133.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.134.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.134.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.134.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.135.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.135.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.135.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.136.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.136.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.136.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.137.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.137.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.137.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.138.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.138.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.138.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.139.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.139.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.139.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.140.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.140.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.140.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.141.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.141.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.141.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.142.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.142.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.142.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.143.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.143.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.143.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.144.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.144.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.144.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.145.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.145.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.145.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.146.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.146.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.146.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.147.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.147.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.147.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.148.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.148.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.148.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.149.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.149.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.149.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.150.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.150.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.150.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.151.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.151.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.151.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.152.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.152.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.152.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.153.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.153.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.153.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.154.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.154.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.154.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.155.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.155.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.155.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.156.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.156.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.156.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.157.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.157.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.157.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.158.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.158.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.158.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.159.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.159.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.159.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.160.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.160.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.160.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.161.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.161.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.161.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.162.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.162.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.162.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.163.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.163.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.163.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.164.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.164.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.164.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.165.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.165.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.165.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.166.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.166.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.166.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.167.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.167.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.167.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.168.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.168.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.168.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.169.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.169.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.169.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.170.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.170.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.170.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.171.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.171.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.171.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.172.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.172.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.172.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.173.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.173.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.173.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.174.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.174.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.174.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.175.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.175.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.175.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.176.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.176.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.176.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.177.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.177.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.177.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.178.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.178.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.178.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.179.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.179.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.179.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.180.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.180.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.180.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.181.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.181.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.181.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.182.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.182.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.182.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.183.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.183.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.183.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.184.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.184.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.184.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.185.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.185.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.185.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.186.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.186.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.186.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.187.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.187.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.187.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.188.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.188.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.188.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.189.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.189.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.189.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.190.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.190.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.190.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.191.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.191.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.191.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.192.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.192.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.192.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.193.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.193.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.193.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.194.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.194.up_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.194.down_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.195.gate_proj.weight": "model-00162-of-000163.safetensors", + "model.layers.61.mlp.experts.195.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.195.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.196.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.196.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.196.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.197.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.197.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.197.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.198.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.198.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.198.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.199.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.199.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.199.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.200.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.200.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.200.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.201.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.201.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.201.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.202.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.202.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.202.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.203.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.203.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.203.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.204.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.204.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.204.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.205.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.205.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.205.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.206.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.206.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.206.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.207.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.207.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.207.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.208.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.208.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.208.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.209.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.209.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.209.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.210.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.210.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.210.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.211.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.211.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.211.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.212.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.212.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.212.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.213.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.213.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.213.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.214.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.214.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.214.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.215.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.215.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.215.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.216.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.216.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.216.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.217.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.217.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.217.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.218.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.218.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.218.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.219.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.219.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.219.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.220.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.220.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.220.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.221.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.221.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.221.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.222.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.222.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.222.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.223.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.223.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.223.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.224.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.224.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.224.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.225.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.225.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.225.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.226.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.226.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.226.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.227.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.227.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.227.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.228.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.228.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.228.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.229.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.229.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.229.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.230.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.230.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.230.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.231.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.231.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.231.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.232.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.232.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.232.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.233.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.233.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.233.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.234.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.234.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.234.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.235.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.235.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.235.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.236.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.236.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.236.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.237.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.237.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.237.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.238.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.238.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.238.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.239.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.239.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.239.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.240.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.240.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.240.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.241.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.241.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.241.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.242.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.242.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.242.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.243.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.243.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.243.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.244.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.244.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.244.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.245.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.245.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.245.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.246.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.246.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.246.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.247.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.247.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.247.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.248.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.248.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.248.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.249.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.249.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.249.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.250.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.250.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.250.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.251.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.251.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.251.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.252.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.252.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.252.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.253.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.253.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.253.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.254.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.254.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.254.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.255.gate_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.255.up_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.mlp.experts.255.down_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.input_layernorm.weight": "model-00163-of-000163.safetensors", + "model.layers.61.post_attention_layernorm.weight": "model-00163-of-000163.safetensors", + "model.layers.61.embed_tokens.weight": "model-00163-of-000163.safetensors", + "model.layers.61.enorm.weight": "model-00163-of-000163.safetensors", + "model.layers.61.hnorm.weight": "model-00163-of-000163.safetensors", + "model.layers.61.eh_proj.weight": "model-00163-of-000163.safetensors", + "model.layers.61.shared_head.norm.weight": "model-00163-of-000163.safetensors", + "model.layers.61.shared_head.head.weight": "model-00163-of-000163.safetensors" + } +} \ No newline at end of file