%------------------------------------------------------------------------------
% File : Vampire---5.0.1
% Problem : CAT019-3 : TPTP v9.3.1. Released v1.0.0.
% Transfm : none
% Format : tptp:raw
% Command : run_vampire /export/starexec/sandbox/benchmark/theBenchmark.p 300 THM
% Computer : n018.cluster.edu
% Model : x86_64 x86_64
% CPU : Intel(R) Xeon(R) CPU E5-2620 v4 2.10GHz
% Memory : 8046.5625MB
% OS : Linux 6.8.0-71-generic
% CPULimit : 300s
% WCLimit : 300s
% DateTime : Tue Sep 29 09:36:39 AM UTC 2026
% Result : Unsatisfiable 3.53s 1.27s
% Output : Refutation 3.53s
% Verified :
% SZS Type : Refutation
% Derivation depth : 13
% Number of leaves : 8
% Syntax : Number of formulae : 33 ( 7 unt; 3 def)
% Number of atoms : 75 ( 26 equ)
% Maximal formula atoms : 4 ( 2 avg)
% Number of connectives : 59 ( 17 ~; 39 |; 0 &)
% ( 3 <=>; 0 =>; 0 <=; 0 <~>)
% Maximal formula depth : 6 ( 3 avg)
% Maximal term depth : 2 ( 1 avg)
% Number of predicates : 6 ( 4 usr; 4 prp; 0-2 aty)
% Number of functors : 3 ( 3 usr; 2 con; 0-2 aty)
% Number of variables : 17 ( 0 sgn 17 !; 0 ?)
% Comments :
%------------------------------------------------------------------------------
fof(f15,axiom,
! [X0,X1] :
( there_exists(f1(X0,X1))
| X0 = X1 ),
file('/export/starexec/sandbox/benchmark/theBenchmark.p',indiscernibles1) ).
fof(f16,axiom,
! [X0,X1] :
( X0 = f1(X0,X1)
| X1 = f1(X0,X1)
| X0 = X1 ),
file('/export/starexec/sandbox/benchmark/theBenchmark.p',indiscernibles2) ).
fof(f17,plain,
! [X0,X1] :
( f1(X0,X1) = X1
| f1(X0,X1) = X0
| X0 = X1 ),
inference(reorient_equations,[],[f16]) ).
fof(f20,axiom,
! [X0] :
( ~ there_exists(X0)
| a != X0
| b = X0 ),
file('/export/starexec/sandbox/benchmark/theBenchmark.p',equality_of_a_and_b1) ).
fof(f21,axiom,
! [X0] :
( ~ there_exists(X0)
| a = X0
| b != X0 ),
file('/export/starexec/sandbox/benchmark/theBenchmark.p',equality_of_a_and_b2) ).
fof(f22,negated_conjecture,
a != b,
file('/export/starexec/sandbox/benchmark/theBenchmark.p',prove_a_equals_b) ).
fof(f25,plain,
( ~ there_exists(a)
| a = b ),
inference(equality_resolution,[],[f20]) ).
fof(f26,plain,
( ~ there_exists(b)
| a = b ),
inference(equality_resolution,[],[f21]) ).
fof(f29,definition,
( spl0_1
<=> a = b ),
introduced(definition,[new_symbols(definition,[spl0_1])],[avatar_definition]) ).
fof(f30,plain,
( a != b
| spl0_1 ),
inference(avatar_component_clause,[],[f29]) ).
fof(f33,definition,
( spl0_2
<=> there_exists(b) ),
introduced(definition,[new_symbols(definition,[spl0_2])],[avatar_definition]) ).
fof(f35,plain,
( ~ there_exists(b)
| spl0_2 ),
inference(avatar_component_clause,[],[f33]) ).
fof(f36,plain,
( spl0_1
| ~ spl0_2 ),
inference(avatar_split_clause,[],[f26,f33,f29]) ).
fof(f38,definition,
( spl0_3
<=> there_exists(a) ),
introduced(definition,[new_symbols(definition,[spl0_3])],[avatar_definition]) ).
fof(f40,plain,
( ~ there_exists(a)
| spl0_3 ),
inference(avatar_component_clause,[],[f38]) ).
fof(f41,plain,
( spl0_1
| ~ spl0_3 ),
inference(avatar_split_clause,[],[f25,f38,f29]) ).
fof(f42,plain,
~ spl0_1,
inference(avatar_split_clause,[],[f22,f29]) ).
fof(f63,plain,
! [X0,X1] :
( there_exists(X0)
| X0 = X1
| f1(X1,X0) = X1
| X0 = X1 ),
inference(superposition,[],[f15,f17]) ).
fof(f65,plain,
! [X0,X1] :
( f1(X1,X0) = X1
| X0 = X1
| there_exists(X0) ),
inference(duplicate_literal_removal,[],[f63]) ).
fof(f127,plain,
! [X0,X1] :
( there_exists(X0)
| X0 = X1
| X0 = X1
| there_exists(X1) ),
inference(superposition,[],[f15,f65]) ).
fof(f128,plain,
! [X0,X1] :
( there_exists(X0)
| X0 = X1
| there_exists(X1) ),
inference(duplicate_literal_removal,[],[f127]) ).
fof(f180,plain,
( ! [X0] :
( there_exists(X0)
| a = X0 )
| spl0_3 ),
inference(resolution,[],[f128,f40]) ).
fof(f190,plain,
( a = b
| spl0_2
| spl0_3 ),
inference(resolution,[],[f180,f35]) ).
fof(f191,plain,
( $false
| spl0_1
| spl0_2
| spl0_3 ),
inference(forward_subsumption_resolution,[],[f190,f30]) ).
fof(f192,plain,
( spl0_1
| spl0_2
| spl0_3 ),
inference(avatar_contradiction_clause,[],[f191]) ).
cnf(s1,plain,
( spl0_1
| ~ spl0_2 ),
inference(sat_conversion,[],[f36]) ).
cnf(s2,plain,
( spl0_1
| ~ spl0_3 ),
inference(sat_conversion,[],[f41]) ).
cnf(s3,plain,
~ spl0_1,
inference(sat_conversion,[],[f42]) ).
cnf(s4,plain,
( spl0_1
| spl0_2
| spl0_3 ),
inference(sat_conversion,[],[f192]) ).
cnf(s5,plain,
~ spl0_3,
inference(rat,[],[s2,s3]) ).
cnf(s6,plain,
spl0_2,
inference(rat,[],[s4,s3,s5]) ).
cnf(s7,plain,
$false,
inference(rat,[],[s1,s6,s3]) ).
fof(f193,plain,
$false,
inference(avatar_sat_refutation,[],[s7]) ).
%------------------------------------------------------------------------------
%----ORIGINAL SYSTEM OUTPUT
% 0.00/0.03 % Problem : CAT019-3 : TPTP v9.3.1. Released v1.0.0.
% 0.00/0.06 % Command : run_vampire /export/starexec/sandbox/benchmark/theBenchmark.p 300 THM
% 0.09/0.19 % Computer : n018.cluster.edu
% 0.09/0.19 % Model : x86_64 x86_64
% 0.09/0.19 % CPU : Intel(R) Xeon(R) CPU E5-2620 v4 @ 2.10GHz
% 0.09/0.19 % Memory : 8046.5625MB
% 0.09/0.19 % OS : Linux 6.8.0-71-generic
% 0.09/0.19 % CPULimit : 300
% 0.09/0.19 % WCLimit : 300
% 0.09/0.19 % DateTime : Mon Sep 28 21:16:11 UTC 2026
% 0.09/0.20 % CPUTime :
% 0.09/0.20 Running run_vampire /export/starexec/sandbox/benchmark/theBenchmark.p 300 THM
% 0.09/0.22 Running first-order theorem proving
% 0.09/0.22 Running: /export/starexec/sandbox/solver/bin/vampire --input_syntax tptp --output_axiom_names on --mode casc -m 16384 --cores 7 -t 300 /export/starexec/sandbox/benchmark/theBenchmark.p
% 3.53/1.27 % (3797457)Input is clausal, will run a generic CNF schedule.
% 3.53/1.27 % (3797466)dis-1002_1_to=lpo:sil=16000:fd=off:random_seed=2224091109:st=1.5:i=114:aac=none:ins=7:ss=axioms:fsd=on_2999 on theBenchmark for (2999ds/114Mi)
% 3.53/1.27 % (3797466)Refutation not found, incomplete strategy
% 3.53/1.27 % (3797466)------------------------------
% 3.53/1.27 % (3797466)Version: Vampire 5.0.1 (Release build, commit ea8961452 on 2026-07-16 15:14:34 +0200)
% 3.53/1.27 % (3797466)Linked with Z3 4.14.0.0 3c47fd96cf5645d0c42b2c819d9e9a84380aa721 z3-4.8.4-9178-g3c47fd96c
% 3.53/1.27 % (3797466)CaDiCaL version: 2.1.3
% 3.53/1.27 % (3797466)Termination reason: Refutation not found, incomplete strategy
% 3.53/1.27 % (3797466)Time elapsed: 0.0000 s
% 3.53/1.27 % (3797466)Peak memory usage: 87 MB
% 3.53/1.27 % (3797468)dis-21_1_sil=8000:lcm=predicate:random_seed=3736048848:st=5:avsq=on:i=117:avsqr=1,16:sd=3:aac=none:ep=RS:fsr=off:ss=included_2999 on theBenchmark for (2999ds/117Mi)
% 3.53/1.27 % (3797463)lrs+10_1_ncem=casc2026/models/loop8.pt:sil=128000:npcc=on:urr=on:br=off:random_seed=251759682:i=132376:av=off_2999 on theBenchmark for (2999ds/132376Mi)
% 3.53/1.27 % (3797462)lrs+10_1_ncem=casc2026/models/loop8.pt:sil=128000:tgt=full:npcc=on:drc=off:sp=weighted_frequency:spb=goal:fd=preordered:foolp=on:random_seed=4137505333:i=140167_2999 on theBenchmark for (2999ds/140167Mi)
% 3.53/1.27 % (3797465)lrs+10_1_sil=8000:sp=occurrence:random_seed=3486055660:i=107:sd=3:ss=axioms:sgt=8_2999 on theBenchmark for (2999ds/107Mi)
% 3.53/1.27 % (3797464)lrs+1002_1_ncem=casc2026/models/loop8.pt:sil=128000:tgt=ground:npcc=on:sp=reverse_frequency:spb=intro:random_seed=3543373683:i=137899:s2at=10:gtgl=3:kws=precedence:add=on:bd=preordered:gtg=position_2999 on theBenchmark for (2999ds/137899Mi)
% 3.53/1.27 % (3797467)dis-1011_1_sil=16000:fde=unused:s2agt=70:random_seed=1055356236:s2a=on:i=180:gtg=position_2999 on theBenchmark for (2999ds/180Mi)
% 3.53/1.27 % (3797465)Refutation not found, incomplete strategy
% 3.53/1.27 % (3797465)------------------------------
% 3.53/1.27 % (3797465)Version: Vampire 5.0.1 (Release build, commit ea8961452 on 2026-07-16 15:14:34 +0200)
% 3.53/1.27 % (3797465)Linked with Z3 4.14.0.0 3c47fd96cf5645d0c42b2c819d9e9a84380aa721 z3-4.8.4-9178-g3c47fd96c
% 3.53/1.27 % (3797465)CaDiCaL version: 2.1.3
% 3.53/1.27 % (3797465)Termination reason: Refutation not found, incomplete strategy
% 3.53/1.27 % (3797465)Time elapsed: 0.001 s
% 3.53/1.27 % (3797465)Peak memory usage: 87 MB
% 3.53/1.27 % (3797467)First to succeed.
% 3.53/1.27 % (3797467)Solution written to "/export/starexec/sandbox/tmp/vampire-proof-3797457"
% 3.53/1.27 % (3797468)Instruction limit reached!
% 3.53/1.27 % (3797468)------------------------------
% 3.53/1.27 % (3797468)Version: Vampire 5.0.1 (Release build, commit ea8961452 on 2026-07-16 15:14:34 +0200)
% 3.53/1.27 % (3797468)Linked with Z3 4.14.0.0 3c47fd96cf5645d0c42b2c819d9e9a84380aa721 z3-4.8.4-9178-g3c47fd96c
% 3.53/1.27 % (3797468)CaDiCaL version: 2.1.3
% 3.53/1.27 % (3797468)Termination reason: Instruction limit
% 3.53/1.27 % (3797468)Termination phase: Saturation
% 3.53/1.27 % (3797468)Time elapsed: 0.057 s
% 3.53/1.27 % (3797468)Peak memory usage: 90 MB
% 3.53/1.27 % (3797468)Instructions burned: 118 (million)
% 3.53/1.27 % (3797466)------------------------------
% 3.53/1.27 % (3797466)------------------------------
% 3.53/1.27 % (3797476)dis+1010_3_sil=8000:plsq=on:drc=off:fde=none:plsqc=1:bsd=on:plsqr=7,2:sos=on:spb=goal_then_units:random_seed=3468837058:i=143:sd=2:aac=none:ss=axioms:sgt=16_2997 on theBenchmark for (2997ds/143Mi)
% 3.53/1.27 % (3797476)Refutation not found, incomplete strategy
% 3.53/1.27 % (3797476)------------------------------
% 3.53/1.27 % (3797476)Version: Vampire 5.0.1 (Release build, commit ea8961452 on 2026-07-16 15:14:34 +0200)
% 3.53/1.27 % (3797476)Linked with Z3 4.14.0.0 3c47fd96cf5645d0c42b2c819d9e9a84380aa721 z3-4.8.4-9178-g3c47fd96c
% 3.53/1.27 % (3797476)CaDiCaL version: 2.1.3
% 3.53/1.27 % (3797476)Termination reason: Refutation not found, incomplete strategy
% 3.53/1.27 % (3797476)Time elapsed: 0.002 s
% 3.53/1.27 % (3797476)Peak memory usage: 88 MB
% 3.53/1.27 % (3797476)Instructions burned: 1 (million)
% 3.53/1.27 % (3797477)ott-1010_1_to=lpo:sil=16000:sos=on:spb=units:urr=on:bce=on:br=off:random_seed=1150601117:st=3:avsq=on:s2a=on:i=189:s2at=1.2:avsqr=1,16:sd=2:bd=all:nm=64:ss=axioms:sgt=30_2997 on theBenchmark for (2997ds/189Mi)
% 3.53/1.27 % (3797477)Also succeeded, but the first one will report.
% 3.53/1.27 % (3797465)------------------------------
% 3.53/1.27 % (3797465)------------------------------
% 3.53/1.27 % (3797467)Refutation found. Thanks to Tanya!
% 3.53/1.27 % SZS status Unsatisfiable for theBenchmark
% 3.53/1.27 % SZS output start Proof for theBenchmark
% See solution above
% 3.53/1.27 % (3797467)------------------------------
% 3.53/1.27 % (3797467)Version: Vampire 5.0.1 (Release build, commit ea8961452 on 2026-07-16 15:14:34 +0200)
% 3.53/1.27 % (3797467)Linked with Z3 4.14.0.0 3c47fd96cf5645d0c42b2c819d9e9a84380aa721 z3-4.8.4-9178-g3c47fd96c
% 3.53/1.27 % (3797467)CaDiCaL version: 2.1.3
% 3.53/1.27 % (3797467)Termination reason: Refutation
% 3.53/1.27 % (3797467)Time elapsed: 0.006 s
% 3.53/1.27 % (3797467)Peak memory usage: 89 MB
% 3.53/1.27 % (3797467)Instructions burned: 7 (million)
% 3.53/1.27 % (3797467)------------------------------
% 3.53/1.27 % (3797467)------------------------------
% 3.53/1.27 % (3797457)Success in time 0.413 s
% 3.53/1.27 % Vampire exiting
%------------------------------------------------------------------------------